Build a network to classify Reuters newswires into 46 mutually exclusive topics.

In [43]:
from keras.datasets import reuters
import numpy as np
from keras import models
from keras import layers
import matplotlib.pyplot as py

################# 3.12 Loading the Reuters dataset #######################################

(train_data, train_labels), (test_data, test_labels) = reuters.load_data(num_words=10000)

# len(train_data)
# train_data[10]
# train_labels
# # len(test_data)

################# 3.13 Decoding newswires back to text #######################################

word_index = reuters.get_word_index()
reverse_word_index = dict([(value, key) for (key, value) in word_index.items()])
decoded_newswire = ''.join([reverse_word_index.get(i - 3, '?') for i in train_data[0]])

# train_data[0]
# train_data[10]

################# 3.14 Encoding the data #####################################################

def vactorize_sequences(sequences, dimension=10000):
    results = np.zeros((len(sequences), dimension)) ### Creates an all-zero matrix of shape (len(sequences), dimension)
    for i, sequence in enumerate(sequences):
        results[i, sequence] = 1. ### Sets specific indices of results[i] to 1s
    return results

x_train = vactorize_sequences(train_data) ### Vectorized training data
x_test = vactorize_sequences(test_data) ### Vectorized test data

############## Vectorize labels us One-hot encoding #################

def to_one_hot(labels, dimension=46):
    results = np.zeros((len(labels), dimension))
    for i, label in enumerate(labels):
        results[i, label] = 1.
    return results

one_hot_train_labels = to_one_hot(train_labels)
one_hot_test_labels = to_one_hot(test_labels)

####################### 3.15 The model definition ######################################

model = models.Sequential()
model.add(layers.Dense(64, activation='relu', input_shape=(10000,)))
model.add(layers.Dense(64, activation='relu'))
model.add(layers.Dense(46, activation='softmax'))

####################### 3.16 Compiling the model ######################################

model.compile(optimizer='rmsprop', 
              loss='categorical_crossentropy', 
              metrics=['accuracy'])

####################### 3.17 Setting aside a validation set #############################

x_val = x_train[:1000]
partial_x_train = x_train[1000:]

y_val = one_hot_train_labels[:1000]
partial_y_train = one_hot_train_labels[1000:]

####################### 3.18 Train the model  ######################################

# history = model.fit(partial_x_train, 
#                     partial_y_train, 
#                     epochs=20, 
#                     batch_size=512, 
#                     validation_data=(x_val, y_val))

####################### 3.19 Plotting the training and validation loss ###############

# loss = history.history['loss']
# val_loss = history.history['val_loss']
# epochs = range(1, len(loss)+1)

# py.plot(epochs, loss, 'bo', label='Training loss')
# py.plot(epochs, val_loss, 'b', label='Validation loss')
# py.title('Training and Validation loss')
# py.xlabel('Epochs')
# py.ylabel('Loss')
# py.legend()
# py.show()

####################### 3.20 Plotting the training and validation Accuracy ###############

# py.clf()

# acc = history.history['acc']
# val_acc = history.history['val_acc']

# py.plot(epochs, loss, 'bo', label='Training Accuracy')
# py.plot(epochs, val_loss, 'b', label='Validation Accuracy')
# py.title('Training and Validation Accuracy')
# py.xlabel('Epochs')
# py.ylabel('Loss')
# py.legend()
# py.show()

####################### 3.21 Retraining the model #########################################

# model.fit(partial_x_train, partial_y_train, epochs=9, validation_data=(x_val, y_val))
# results = model.evaluate(x_test, one_hot_test_labels)
# results ### Accuracy = 78.85%

####################### 3.22 generating prediction for new data ############################

prediction = model.predict(x_test)
prediction[0].shape ### Each entry in predictions is a vector of length 46
np.sum(prediction[0]) ### The coefficients in this vector sum to 1
np.argmax(prediction[1]) ### The largest entry is the predicted class with the highest probability
# prediction[0]

42