In [1]:
from sklearn.svm import SVC
from sklearn.neural_network import MLPClassifier
from sklearn.naive_bayes import BernoulliNB
from sklearn.model_selection import cross_val_score
from py.utils import load_data

directory = '../data/'
heads = ['l30_r15', 'l10_r10', 'l5_r5']
n_cv = 5

In [2]:
import pickle
performances = {}

for head in heads:
    print('\n\nhead = %s' % head)
    x, y, x_words, vocabs = load_data(head, directory)
    
    classifier = BernoulliNB()
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('\nBernoulli Naive Bayes: ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('BernoulliNB', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'BernoulliNB ' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   
        
    
    classifier = MLPClassifier(hidden_layer_sizes=(5,))
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Multilayer Perceptron Classifier (h=[5]): ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('MLPClassifier (5,)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'MLPClassifier (5,)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    classifier = MLPClassifier(hidden_layer_sizes=(20,))
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Multilayer Perceptron Classifier (h=[20])', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('MLPClassifier (20,)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'MLPClassifier (20,)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    classifier = MLPClassifier(hidden_layer_sizes=(50,10))
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Multilayer Perceptron Classifier (h=[50, 10]): ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('MLPClassifier (50,10)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'MLPClassifier (50,10)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    classifier = SVC(C=10.0, kernel='rbf',shrinking=True)
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Support Vector Machine (rbf, C=10.0): ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('SVC (C=10)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = ' SVC (rbf, C=10.0)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    classifier = SVC(C=1.0, kernel='rbf',shrinking=True)
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Support Vector Machine (rbf, C=1.0): ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('SVC (rbf, C=1.0)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'SVC (rbf, C=1.0)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    classifier = SVC(C=0.1, kernel='rbf',shrinking=True)
    scores = cross_val_score(classifier, x, y, cv=n_cv)
    print('Support Vector Machine (rbf, C=0.1): ', end='')
    print(' > %s' % ['%.5f' % s for s in scores])
    performances[('SVC (rbf, C=0.1)', head)] = scores
    with open('performance_other_classifier.pkl', 'wb') as f:
        pickle.dump(performances, f)
    classifier.fit(x, y)
    model_name = 'SVC (C=0.1)' + head
    with open('../models/%s.pkl' % model_name, 'wb') as f:
        pickle.dump(classifier, f)   

    
    print('-' * 80)



head = l30_r15
x shape = (15166, 2617)
y shape = (15166,)
# features = 2617
# L words = 15166

Bernoulli Naive Bayes:  > ['0.99011', '0.98681', '0.98648', '0.98615', '0.98582']
Multilayer Perceptron Classifier (h=[5]):  > ['0.99242', '0.98945', '0.99077', '0.99143', '0.98879']
Multilayer Perceptron Classifier (h=[20]) > ['0.99242', '0.98978', '0.98912', '0.99077', '0.98879']
Multilayer Perceptron Classifier (h=[50, 10]):  > ['0.99242', '0.98945', '0.98978', '0.99110', '0.98945']
Support Vector Machine (rbf, C=10.0):  > ['0.96407', '0.95846', '0.96505', '0.96934', '0.96274']
Support Vector Machine (rbf, C=1.0):  > ['0.95946', '0.95780', '0.95978', '0.96604', '0.96274']
Support Vector Machine (rbf, C=0.1):  > ['0.93507', '0.93867', '0.94362', '0.94362', '0.94230']
--------------------------------------------------------------------------------


head = l10_r10
x shape = (31797, 3297)
y shape = (31797,)
# features = 3297
# L words = 31797

Bernoulli Naive Bayes:  > ['0.98145', '0.98302'