In [2]:
from google.colab import drive
drive.mount('/content/drive')

Drive already mounted at /content/drive; to attempt to forcibly remount, call drive.mount("/content/drive", force_remount=True).


In [3]:
# LIBRARIES
import pickle
import nltk
from textblob import TextBlob
from nltk.corpus import stopwords
from nltk.stem import PorterStemmer
import re

# Install dependencies if not already installed
!pip install -q textblob nltk

# Load Pickle Files
from google.colab import files
uploaded = files.upload()  # Upload your 'best_model.pkl' and 'count_vectorizer.pkl'

model = pickle.load(open('best_model.pkl', 'rb'))
vectorizer = pickle.load(open('count_vectorizer.pkl', 'rb'))

# Download NLTK Data
nltk.download('stopwords')

# TEXT PREPROCESSING
sw = set(stopwords.words('english'))
def text_preprocessing(text):
    txt = TextBlob(text)
    result = txt.correct()
    removed_special_characters = re.sub("[^a-zA-Z]", " ", str(result))
    tokens = removed_special_characters.lower().split()
    stemmer = PorterStemmer()

    cleaned = []
    stemmed = []

    for token in tokens:
        if token not in sw:
            cleaned.append(token)

    for token in cleaned:
        token = stemmer.stem(token)
        stemmed.append(token)

    return " ".join(stemmed)

# TEXT CLASSIFICATION
def text_classification(text):
    if len(text) < 1:
        print("Please enter a review!")
    else:
        print("Classification in progress...")
        cleaned_review = text_preprocessing(text)
        process = vectorizer.transform([cleaned_review]).toarray()
        prediction = model.predict(process)
        p = ''.join(str(i) for i in prediction)

        if p == 'True':
            print("The review entered is Legitimate.")
        elif p == 'False':
            print("The review entered is Fraudulent.")

# Main Function
def main():
    print("Fraud Detection in Online Consumer Reviews Using Machine Learning Techniques\n")

    print("[INFO] Model: Logistic Regression")
    print("[INFO] Vectorizer: Count Vectorizer\n")

    # Input review from the user
    review = input("Enter your review: ")
    text_classification(review)

# Run Main
if __name__ == "__main__":
    main()


https://scikit-learn.org/stable/model_persistence.html#security-maintainability-limitations
https://scikit-learn.org/stable/model_persistence.html#security-maintainability-limitations
[nltk_data] Downloading package stopwords to /root/nltk_data...
[nltk_data]   Unzipping corpora/stopwords.zip.


Saving count_vectorizer.pkl to count_vectorizer.pkl
Saving best_model.pkl to best_model.pkl
Fraud Detection in Online Consumer Reviews Using Machine Learning Techniques

[INFO] Model: Logistic Regression
[INFO] Vectorizer: Count Vectorizer

Enter your review: best prodect
Classification in progress...
The review entered is Legitimate.
