In [1]:
# Import modules
import numpy as np
import matplotlib.pyplot as plt
from sklearn import model_selection

# Import PySwarms
import pyswarms as ps

# Some more magic so that the notebook will reload external python modules;
# see http://stackoverflow.com/questions/1907993/autoreload-of-modules-in-ipython
%reload_ext autoreload
%autoreload 2

In [2]:
# Load the dataset
import pandas as pd
data = pd.read_csv('parkinsons.csv', delimiter=',')
data.columns

Index(['name', 'MDVP:Fo(Hz)', 'MDVP:Fhi(Hz)', 'MDVP:Flo(Hz)', 'MDVP:Jitter(%)',
       'MDVP:Jitter(Abs)', 'MDVP:RAP', 'MDVP:PPQ', 'Jitter:DDP',
       'MDVP:Shimmer', 'MDVP:Shimmer(dB)', 'Shimmer:APQ3', 'Shimmer:APQ5',
       'MDVP:APQ', 'Shimmer:DDA', 'NHR', 'HNR', 'status', 'RPDE', 'DFA',
       'spread1', 'spread2', 'D2', 'PPE'],
      dtype='object')

In [12]:
# Store the features as X and the labels as y
X = data.drop(columns=['name', 'status']).to_numpy()
y = data['status'].to_numpy()
data.columns

# four best features as per mid sems
X = data[['HNR','RPDE','DFA','PPE']].to_numpy()
X.shape

# set 1 as per correlation coeff -> 83.58% training, 78.46% testing
X = data[['MDVP:Jitter(Abs)','MDVP:RAP','MDVP:PPQ','Jitter:DDP']].to_numpy()
X.shape

# set 2 as per correlation coeff -> 83.58% training, 80% testing
X = data[['MDVP:Jitter(Abs)','MDVP:RAP','Jitter:DDP']].to_numpy()
X.shape

(195, 4)

In [13]:
# 60% training set and 40% testing set
x_train, x_test, y_train, y_test = model_selection.train_test_split (X, y, test_size=0.33, random_state=0)
x_train.shape

(130, 4)

In [5]:
X.dtype

dtype('float64')

In [6]:
def sigmoid(Z):
    return 1/(1+np.exp(-Z))

def relu(Z):
    return np.maximum(0,Z)

def sigmoid_backward(dA, Z):
    sig = sigmoid(Z)
    return dA * sig * (1 - sig)

def relu_backward(dA, Z):
    dZ = np.array(dA, copy = True)
    dZ[Z <= 0] = 0;
    return dZ;

In [7]:
# Forward propagation
def forward_prop(params):
    """Forward propagation as objective function

    This computes for the forward propagation of the neural network, as
    well as the loss. It receives a set of parameters that must be
    rolled-back into the corresponding weights and biases.

    Inputs
    ------
    params: np.ndarray
        The dimensions should include an unrolled version of the
        weights and biases.

    Returns
    -------
    float
        The computed negative log-likelihood loss given the parameters
    """
    # Neural network architecture
    n_inputs = 3
    n_hidden = 20
    n_classes = 2

    # Roll-back the weights and biases
    W1 = params[0:60].reshape((n_inputs,n_hidden))
    b1 = params[60:80].reshape((n_hidden,))
    W2 = params[80:120].reshape((n_hidden,n_classes))
    b2 = params[120:122].reshape((n_classes,))

    # Perform forward propagation
    z1 = X.dot(W1) + b1  # Pre-activation in Layer 1
    a1 = np.tanh(z1)     # Activation in Layer 1
    z2 = a1.dot(W2) + b2 # Pre-activation in Layer 2
    logits = z2          # Logits for Layer 2

    # Compute for the softmax of the logits
    exp_scores = np.exp(logits)
    probs = exp_scores / np.sum(exp_scores, axis=1, keepdims=True)

    # Compute for the negative log likelihood
    N = 195 # Number of samples
    corect_logprobs = -np.log(probs[range(N), y])
    loss = np.sum(corect_logprobs) / N

    return loss

In [8]:
def f(x):
    """Higher-level method to do forward_prop in the
    whole swarm.

    Inputs
    ------
    x: numpy.ndarray of shape (n_particles, dimensions)
        The swarm that will perform the search

    Returns
    -------
    numpy.ndarray of shape (n_particles, )
        The computed loss for each particle
    """
    n_particles = x.shape[0]
    j = [forward_prop(x[i]) for i in range(n_particles)]
    return np.array(j)

In [9]:
%%time
# Initialize swarm
options = {'c1': 0.5, 'c2': 0.5, 'w':0.9}

# Call instance of PSO
dimensions = (3 * 20) + (20 * 2) + 20 + 2
optimizer = ps.single.GlobalBestPSO(n_particles=100, dimensions=dimensions, options=options)

# Perform optimization
cost, pos = optimizer.optimize(f, iters=1000)

2019-12-09 17:07:03,448 - pyswarms.single.global_best - INFO - Optimize for 1000 iters with {'c1': 0.5, 'c2': 0.5, 'w': 0.9}
pyswarms.single.global_best: 100%|██████████|1000/1000, best_cost=0.425
2019-12-09 17:07:28,471 - pyswarms.single.global_best - INFO - Optimization finished | best cost: 0.424905281779306, best pos: [ 1.28454854e+01  5.54996491e+00  3.57847084e+00  2.21827300e+00
  4.22990235e+00  7.05954447e-01  4.12499798e+00  3.06478051e-01
  3.86884631e+00  5.11272251e-01  1.40104442e+01  4.23900300e+00
 -1.14137852e+00  1.67221640e-01  1.15101716e+05  6.73172704e+00
 -3.05278086e+01 -1.05251167e-01  2.25815523e+00 -8.14175330e-01
  3.72058958e-01  1.13228700e+00 -1.54559385e+00  2.81431414e+00
 -8.97798494e-01 -6.11515543e-01  1.44623037e+00  5.70337065e-01
  1.18904815e+01 -8.74235211e-01  2.40457562e+00 -8.31352167e-01
  4.49957403e+00  1.56336503e+00 -1.56236229e+00 -2.16782083e+00
  3.51146760e+00  8.62261317e-01  2.93109907e+00 -3.92472044e-01
  1.93010184e+00 -8.711782

CPU times: user 25.1 s, sys: 364 ms, total: 25.4 s
Wall time: 25 s


In [10]:
def predict(X, pos):
    """
    Use the trained weights to perform class predictions.

    Inputs
    ------
    X: numpy.ndarray
        Input Iris dataset
    pos: numpy.ndarray
        Position matrix found by the swarm. Will be rolled
        into weights and biases.
    """
    # Neural network architecture
    n_inputs = 3
    n_hidden = 20
    n_classes = 2

    # Roll-back the weights and biases
    W1 = pos[0:60].reshape((n_inputs,n_hidden))
    b1 = pos[60:80].reshape((n_hidden,))
    W2 = pos[80:120].reshape((n_hidden,n_classes))
    b2 = pos[120:122].reshape((n_classes,))

    # Perform forward propagation
    z1 = X.dot(W1) + b1  # Pre-activation in Layer 1
    a1 = np.tanh(z1)     # Activation in Layer 1
    z2 = a1.dot(W2) + b2 # Pre-activation in Layer 2
    logits = z2          # Logits for Layer 2

    y_pred = np.argmax(logits, axis=1)
    return y_pred

In [11]:
(predict(X, pos) == y).mean()

0.8358974358974359