In [1]:
from sklearn.datasets import fetch_mldata
DATA_PATH = '~/data'
mnist = fetch_mldata('MNIST original', data_home=DATA_PATH)

In [2]:
import matplotlib.pyplot as plt
from sklearn.datasets import fetch_mldata
from sklearn.neural_network import MLPClassifier
import numpy as np
from sklearn.metrics import confusion_matrix
from sklearn.decomposition import PCA
import time

time_start = time.clock()

        
def my_callback(event, **variables):
    print(event) # The name of the event, as shown in the list above.
    print(variables)


mnist = fetch_mldata("MNIST original")
# rescale the data, use the traditional train/test split
X, y = mnist.data / 255., mnist.target    
X_train,X_test = X[:60000], X[60000:]
y_train,y_test = y[:60000], y[60000:]
X_val, y_val = X[50000:60000], y[50000:60000]

# mlp = MLPClassifier(hidden_layer_sizes=(100, 100), max_iter=400, alpha=1e-4,
#                     solver='sgd', verbose=10, tol=1e-4, random_state=1)

'''score = {}
e = PCA(n_components = 100)
e_p = e.fit(X_train)
train = e_p.transform(X_train)
test = e_p.transform(X_test)'''

score = {}
con_m = []
for n in [50,130]:
    mlp = MLPClassifier(hidden_layer_sizes=(n,),activation  = 'relu',
                        solver="lbfgs", learning_rate_init=.001) #lbfgs
    nn = mlp.fit(X_train, y_train)
    score[n] = (nn.score(X_test, y_test), nn.loss_)
    con_m.append(confusion_matrix(y_test,nn.predict(X_test)))

time_elapsed = (time.clock() - time_start)     


In [None]:
import numpy as np 
import math
import random
import matplotlib  
import matplotlib.pyplot as plt 
import sklearn as sklearn
from sklearn.datasets import fetch_mldata
from sklearn.neural_network import MLPClassifier

from sklearn.metrics import confusion_matrix

random.seed(0)

def sigmod_derivate(x):
    return x * (1 - x)
    
def rand(a, b):
    return (b - a) * random.random() + a


def matrix(m, n, fill=0.0):
    mat = []
    for i in range(m):
        mat.append([fill] * n)
    return mat

def sigmoid(x):
    return 1.0 / (1.0 + math.exp(-x))


class BP_NeuralNetwork:
    def __init__(self):
        self.input_n = 0
        self.hidden_n = 0
        self.output_n = 0
        self.input_cells = []
        self.hidden_cells = []
        self.output_cells = []
        self.input_weights = []
        self.output_weights = []
        self.input_correction = []
        self.output_correction = []
        self.MSE = [] 

    def setup(self, ni, nh, no):
        self.input_n = ni + 1
        self.hidden_n = nh
        self.output_n = no
        # init cells
        self.input_cells = [1.0] * self.input_n
        self.hidden_cells = [1.0] * self.hidden_n
        self.output_cells = [1.0] * self.output_n
        # init weights
        self.input_weights = matrix(self.input_n, self.hidden_n)
        self.output_weights = matrix(self.hidden_n, self.output_n)
        # random activate
        for i in range(self.input_n):
            for h in range(self.hidden_n):
                self.input_weights[i][h] = rand(-0.2, 0.2)
        for h in range(self.hidden_n):
            for o in range(self.output_n):
                self.output_weights[h][o] = rand(-2.0, 2.0)
        # init correction matrix
        self.input_correction = matrix(self.input_n, self.hidden_n)
        self.output_correction = matrix(self.hidden_n, self.output_n)

    def predict(self, inputs):
        # activate input layer
        for i in range(self.input_n - 1):
            self.input_cells[i] = inputs[i]
        # activate hidden layer
        for j in range(self.hidden_n):
            total = 0.0
            for i in range(self.input_n):
                total += self.input_cells[i] * self.input_weights[i][j]
            self.hidden_cells[j] = sigmoid(total)
        # activate output layer
        for k in range(self.output_n):
            total = 0.0
            for j in range(self.hidden_n):
                total += self.hidden_cells[j] * self.output_weights[j][k]
            self.output_cells[k] = sigmoid(total)
        return self.output_cells[:]

    def back_propagate(self, case, label, learn, correct):
        # feed forward
        self.predict(case)
        # get output layer error
        output_deltas = [0.0] * self.output_n
        for o in range(self.output_n):
            error = label - self.output_cells[o]
            output_deltas[o] = sigmod_derivate(self.output_cells[o]) * error
        # get hidden layer error
        hidden_deltas = [0.0] * self.hidden_n
        for h in range(self.hidden_n):
            error = 0.0
            for o in range(self.output_n):
                error += output_deltas[o] * self.output_weights[h][o]
            hidden_deltas[h] = sigmod_derivate(self.hidden_cells[h]) * error
        # update output weights
        for h in range(self.hidden_n):
            for o in range(self.output_n):
                change = output_deltas[o] * self.hidden_cells[h]
                self.output_weights[h][o] += learn * change + correct * self.output_correction[h][o]
                self.output_correction[h][o] = change
        # update input weights
        for i in range(self.input_n):
            for h in range(self.hidden_n):
                change = hidden_deltas[h] * self.input_cells[i]
                self.input_weights[i][h] += learn * change + correct * self.input_correction[i][h]
                self.input_correction[i][h] = change
        # get global error
        error = 0.0        
        for o in range(1):
            error += 0.5 * (label - self.output_cells[o]) ** 2
            
        return error

    def train(self, cases, labels, limit=10000, learn=0.05, correct=0.1):
        for i in range(limit):
            error = 0.0
            for i in range(len(cases)):
                label = labels[i]
                case = cases[i]
                error += self.back_propagate(case, label, learn, correct)
                self.MSE.append(error)
        return error,self.MSE

    def test(self):
        mnist = fetch_mldata("MNIST original")
        # rescale the data, use the traditional train/test split
        X, y = mnist.data / 255., mnist.target    
        X_train,X_test = X[:60000], X[60000:]
        y_train,y_test = y[:60000], y[60000:]
        X_val, y_val = X[50000:60000], y[50000:60000]
        #predictions = []
        self.setup(784,15,1)
        err, cost = self.train(X_train, y_train, 10000, 0.05, 0.1)
        return err,cost
        '''
        for point in points:
            print(self.predict(point))
            predictions.append(self.predict(point))
        return points,predictions'''

if __name__ == '__main__':
    nn = BP_NeuralNetwork()
    f_err, MSE_list = nn.test()
    
    