<a href="https://colab.research.google.com/github/HaoYamado/notebooks/blob/master/Feedforward_Neural_Network.ipynb" target="_parent"><img src="https://colab.research.google.com/assets/colab-badge.svg" alt="Open In Colab"/></a>

# Feedforward Neural Network with PyTorch

## About Feedforward Neural Network
### Logistic Regression Transition to Neural Network
#### Logistic Regression Review

In [59]:
# Define logistic regression model
import torch
import torch.nn as nn

# Define our model class
class LogisticRegressionModel(nn.Module):
  def __init__(self, input_dim, output_dim):
    super(LogisticRegressionModel, self).__init__()
    self.linear = nn.Linear(input_dim, output_dim)
  
  def forward(self, x):
    out = self.linear(x)
    return out

# instantiate the logistic regression model
input_dim = 28*28
output_dim = 10

model = LogisticRegressionModel(input_dim, output_dim)
print(model)

LogisticRegressionModel(
  (linear): Linear(in_features=784, out_features=10, bias=True)
)


## Build Feedforward Neural Network with PyTorch

#### Model A: 1 Hidden layer Feedforward Neural Network (Sigmoid activation)
![alt text](https://www.deeplearningwizard.com/deep_learning/practical_pytorch/images/nn1.png)

#### Steps:
*   Step 1: Load Dataset
*   Step 2: Make Dataset Iterable
*   Step 3: Create Model Class
*   Step 4: Instantiate Model Class
*   Step 5: Instantiate Loss Class
*   Step 6: Instantiate Optimizer Class
*   Step 7: Train Model

## Step 1: Loading MNIST Train Dataset



In [0]:
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets

train_dataset = dsets.MNIST(root='./data',
                            train=True,
                            transform=transforms.ToTensor(),
                            download=True)
test_dataset = dsets.MNIST(root='./data',
                          train=False,
                          transform=transforms.ToTensor())

## Step 2: Make Dataset Iterable

In [61]:
# Batch sizes and iterations

60000/100

600.0

In [62]:
 # Epochs
 600.0 * 5

3000.0

In [0]:
# Brinding batch size, iterations and epochs together

batch_size = 100
n_iters = 3000
num_epochs = n_iters / (len(train_dataset) / batch_size)
num_epochs = int(num_epochs)

train_loader = torch.utils.data.DataLoader(dataset=train_dataset,
                                           batch_size = batch_size,
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset,
                                          batch_size = batch_size,
                                          shuffle=False)

## Step 3: Create Model Class

In [0]:
class FeedforwardNeuralNetModel(nn.Module):
  def __init__(self, input_dim, hidden_dim, output_dim):
    super(FeedforwardNeuralNetModel, self).__init__()
    # linear function
    self.fc1 = nn.Linear(input_dim, hidden_dim)

    # Non-linearity
    self.sigmoid = nn.Sigmoid()

    # Linear function (readout)
    self.fc2 = nn.Linear(hidden_dim, output_dim)
  
  def forward(self, x):
    # linear function # LINEAR
    out = self.fc1(x)

    # Non-linearity # NON LINEAR
    out = self.sigmoid(out)

    # Linear function (readout) # LINEAR
    out = self.fc2(out)
    return out

## Step 4: Instantiate Model Class

In [0]:
input_dim = 28*28
hidden_dim = 100
output_dim = 10

model = FeedforwardNeuralNetModel(input_dim, hidden_dim, output_dim)

## Step 5: Instantiate Loss Class

In [0]:
criterion = nn.CrossEntropyLoss()

## Step 6: Instantiate Optimizer Class

In [0]:
learning_rate = 0.1
optimizer = torch.optim.SGD(model.parameters(), lr=learning_rate)

In [68]:
print(model.parameters())
print(len(list(model.parameters())))
# FC 1
print(list(model.parameters())[0].size())
# FC 1 Bieas Parameters
print(list(model.parameters())[1].size())
# FC 2 Parameters
print(list(model.parameters())[2].size())
# FC 2 Bias Parameters
print(list(model.parameters())[3].size())

<generator object Module.parameters at 0x7f8bc9442888>
4
torch.Size([100, 784])
torch.Size([100])
torch.Size([10, 100])
torch.Size([10])


![alt text](https://www.deeplearningwizard.com/deep_learning/practical_pytorch/images/nn1_params3.png)

## Step 7: Train Model
#### Process:
*   a. Convert inputs to tensors with gradients accumulation capabilities
*   b. Clear gradients buffers
*   c. Get output given inputs
*   d. Get loss
*   e. Get gradients w.r.t. parameters
*   f. Update parameters using gradients
    *   parameters = parameters - learning_rate * parameters_gradients
*   g. REPEAT




In [69]:
# training process
iter = 0
for epoch in range(num_epochs):
  for i, (images, labels) in enumerate(train_loader):
    # load images with gradient accumulation capabilities
    images = images.view(-1, 28*28).requires_grad_()
    # clear gradients w.r.t. parameters
    optimizer.zero_grad()
    # forward pass to get output/logits
    outputs = model(images)
    # Calculate Loss: softmax ---> cross entropy loss
    loss = criterion(outputs, labels)
    # getting gradients w.r.t. parameters
    loss.backward()
    # updating parameters
    optimizer.step()
    iter += 1

    if iter % 500 == 0:
      # Calculate Accuracy
      correct = 0
      total = 0
      # iterate through test dataset
      for images, labels in test_loader:
        # load images with gradient accumulation capabilities
        images = images.view(-1, 28*28).requires_grad_()
        # forward pass only to get logits/ouput
        outputs = model(images)
        # get predictions from maximum value
        _, predicted = torch.max(outputs.data, 1)
        #  total number of labels
        total += labels.size(0)
        # Total correct predictions
        correct += (predicted == labels).sum()
      accuracy = 100 * correct/total

      print('Iteration: {}. Loss: {}. Accuracy: {}'.format(iter, loss.item(), accuracy))

Iteration: 500. Loss: 0.6887582540512085. Accuracy: 86
Iteration: 1000. Loss: 0.46380242705345154. Accuracy: 89
Iteration: 1500. Loss: 0.43305614590644836. Accuracy: 90
Iteration: 2000. Loss: 0.2959619462490082. Accuracy: 91
Iteration: 2500. Loss: 0.3037315607070923. Accuracy: 91
Iteration: 3000. Loss: 0.3642258048057556. Accuracy: 92


# Model B: 1 Hidden Layer Feedforward Neural Network (Tanh Activation)

![alt text](https://www.deeplearningwizard.com/deep_learning/practical_pytorch/images/nn1.png)

#### Steps:


*   Step 1: Load Dataset
*   Step 2: Make Dataset iterable
*   Step 3: Create Model Class
*   Step 4: Instantiate Model Class
*   Step 5: Instantiate Loss Class
*   Step 6: Insantiate Optimizer Class
*   Step 7: Train Model


In [70]:
# 1-layer with Tanh activation
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets

"""
STEP 1: LOADING DATASET
"""
train_dataset = dsets.MNIST(root='./data',
                            train=True,
                            transform=transforms.ToTensor(),
                            download=True)

test_dataset = dsets.MNIST(root='./data',
                           train=False,
                           transform=transforms.ToTensor())

"""
STEP 2: MAKING DATASET ITERABLE
"""

batch_size = 100
n_iters = 3000
num_epochs = n_iters / (len(train_dataset) / batch_size)
num_epochs = int(num_epochs)

train_loader = torch.utils.data.DataLoader(dataset=train_dataset,
                                           batch_size=batch_size,
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset,
                                          batch_size=batch_size,
                                          shuffle=False)


"""
STEP 3: CREATE MODEL CLASS
"""
class FeedforwardNeuralNetModel(nn.Module):
  def __init__(self, input_dim, hidden_dim, output_dim):
    super(FeedforwardNeuralNetModel, self).__init__()
    # Linear Function
    self.fc1 = nn.Linear(input_dim, hidden_dim)
    # Non-lenearity
    self.tanh = nn.Tanh()
    # Linear function (readout)
    self.fc2 = nn.Linear(hidden_dim, output_dim)
  
  def forward(self, x):
    # Linear funstion
    out = self.fc1(x)
    # Non-linearity
    out = self.tanh(out)
    # Linear function (readout)
    out = self.fc2(out)
    return out

"""
STEP 4: INSTANTIATE MODEL CLASS
"""
imput_dim = 28*28
hidden_dim = 100
output_dim = 10

model = FeedforwardNeuralNetModel(input_dim, hidden_dim, output_dim)

"""
STEP 5: INSTANTIATE LOSS CLASS
"""
criterion = nn.CrossEntropyLoss()
"""
STEP 6: INSTANTIATE OPTIMIZER CLASS
"""
learning_rate = 0.1
optimizer = torch.optim.SGD(model.parameters(), lr=learning_rate)

"""
STEP 7: TRAIN THE MODEL
"""

iter = 0
for epoch in range(num_epochs):
  for i, (images, labels) in enumerate(train_loader):
    # load images with gradient accumulation capabilities
    images = images.view(-1, 28*28).requires_grad_()
    #Clear gradient w.r.t. parameters
    optimizer.zero_grad()
    # forward pass to get output.logits
    outputs = model(images)
    # calculate Los: softmax ---> cross entropy loss
    loss = criterion(outputs, labels)
    # getting gradients w.r.t. parameters
    loss.backward()
    # Updating paramers
    optimizer.step()

    iter += 1
    if iter % 500 == 0:
      # Calculate Accuracy
      correct = 0
      total = 0
      # iterate through test dataset
      for images, labels in test_loader:
        # load images with gradient accumulation capabilities
        images = images.view(-1, 28*28).requires_grad_()
        # forward pass only to get logits/output
        outputs = model(images)
        # get prediction from maximum value
        _, predicted = torch.max(outputs.data, 1)
        # total number of labels
        total += labels.size(0)
        # total correct predictions
        correct += (predicted == labels).sum()
      accuracy = 100 * correct / total
      # print loss 
      print('Iteration: {}. Loss: {}. Accuracy: {}'.format(iter, loss.item(), accuracy))


Iteration: 500. Loss: 0.49559611082077026. Accuracy: 91
Iteration: 1000. Loss: 0.3960130214691162. Accuracy: 92
Iteration: 1500. Loss: 0.21708333492279053. Accuracy: 93
Iteration: 2000. Loss: 0.1590883880853653. Accuracy: 93
Iteration: 2500. Loss: 0.1721712350845337. Accuracy: 94
Iteration: 3000. Loss: 0.20714856684207916. Accuracy: 95


# Model C: 1 Hidden Layer Feedforward Neural Network (ReLU activation)
![alt text](https://www.deeplearningwizard.com/deep_learning/practical_pytorch/images/nn1.png)

In [71]:
#@title Model C: Hidden layer(Relu activation)
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets

"""
STEP 1: LOADING DATASET
"""

train_datset = dsets.MNIST(root='./data',
                           train=True,
                           transform=transforms.ToTensor(),
                           download=True)

test_dataset = dsets.MNIST(root='./data',
                           train=False,
                           transform=transforms.ToTensor())

"""
STEP 2: MAKING DATASET ITERABLE
"""

batch_size =100
n_iters = 3000
num_epochs = n_iters /(len(train_dataset) / batch_size)
num_epochs = int(num_epochs)

train_loader = torch.utils.data.DataLoader(dataset=train_dataset,
                                           batch_size=batch_size,
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset,
                                          batch_size=batch_size,
                                          shuffle=False)

"""
STEP 3: Create Model Class
"""

class FeedforwardNeuralNetModel(nn.Module):
  def __init__(self, input_dim, hidden_dim, output_dim):
    super(FeedforwardNeuralNetModel, self).__init__()
    # Linear function
    self.fc1 = nn.Linear(input_dim, hidden_dim)
    # Non-Linearity
    self.relu = nn.ReLU()
    # Linear function (readout)
    self.fc2 = nn.Linear(hidden_dim, output_dim)
  
  def forward(self, x):
     # linear function
     out = self.fc1(x)
     # Non-linearity
     out = self.relu(out)
     # Linear function (readout)
     out = self.fc2(out)
     return out
    
"""
STEP 4: INSTANTIATE MODEL CLASS
"""

input_dim = 28*28
hidden_dim = 100
output_dim = 10

model = FeedforwardNeuralNetModel(input_dim, hidden_dim, output_dim)
"""
STEP 5: INSTANTIATE LOSS CLASS
"""
criterion = nn.CrossEntropyLoss()

"""
STEP 6: INSTANTIATE OPTIMIZER CLASS
"""
learning_rate = 0.1
optimizer = torch.optim.SGD(model.parameters(), lr=learning_rate)
"""
STEP 7: TRAIN THE MODEL
"""
iter = 0
for epoch in range(num_epochs):
  for i, (images, labels) in enumerate(train_loader):
    # load images with gradient accumulation capabilities
    images = images.view(-1, 28*28).requires_grad_()
    # Clear gradients w.r.t. parameters
    optimizer.zero_grad()
    # forward pass to get output/logits
    outputs = model(images)
    # Calculate Loss: softmax--->cross entropy loss
    loss = criterion(outputs, labels)
    # Getting gradients w.r.t. parameters
    loss.backward()
    # updating parameters
    optimizer.step()

    iter += 1
    if iter % 500 == 0:
      # Calculate Accuracy
      correct = 0
      total = 0
      # Iterate through test dataset
      for images, labels in test_loader:
        # Load images with gradients accumulation capabilities 
        images = images.view(-1, 28*28).requires_grad_()
        # forward pass only to get logits/output
        outputs = model(images)
        # get predictions form the maximum value
        _, predicted = torch.max(outputs.data, 1)
        # total number of labels 
        total += labels.size(0)
        # total correct predictions
        correct += (predicted == labels).sum()
      accuracy = 100 * correct / total
      # print loss
      print('Iteration: {}. Loss: {}. Accuracy: {}'.format(iter, loss.item(), accuracy))

Iteration: 500. Loss: 0.42539066076278687. Accuracy: 91
Iteration: 1000. Loss: 0.25788912177085876. Accuracy: 92
Iteration: 1500. Loss: 0.15192098915576935. Accuracy: 93
Iteration: 2000. Loss: 0.3578575849533081. Accuracy: 94
Iteration: 2500. Loss: 0.12928007543087006. Accuracy: 95
Iteration: 3000. Loss: 0.17368324100971222. Accuracy: 95


# Model D: 2 Hidden Layer Feedforward Neural Network(ReLU activation)
![alt text](https://www.deeplearningwizard.com/deep_learning/practical_pytorch/images/nn2.png)


In [72]:
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets

'''
STEP 1: LOADING DATASET
'''

train_dataset = dsets.MNIST(root='./data', 
                            train=True, 
                            transform=transforms.ToTensor(),
                            download=True)

test_dataset = dsets.MNIST(root='./data', 
                           train=False, 
                           transform=transforms.ToTensor())

'''
STEP 2: MAKING DATASET ITERABLE
'''

batch_size = 100
n_iters = 3000
num_epochs = n_iters / (len(train_dataset) / batch_size)
num_epochs = int(num_epochs)

train_loader = torch.utils.data.DataLoader(dataset=train_dataset, 
                                           batch_size=batch_size, 
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset, 
                                          batch_size=batch_size, 
                                          shuffle=False)

'''
STEP 3: CREATE MODEL CLASS
'''
class FeedforwardNeuralNetModel(nn.Module):
    def __init__(self, input_dim, hidden_dim, output_dim):
        super(FeedforwardNeuralNetModel, self).__init__()
        # Linear function 1: 784 --> 100
        self.fc1 = nn.Linear(input_dim, hidden_dim) 
        # Non-linearity 1
        self.relu1 = nn.ReLU()

        # Linear function 2: 100 --> 100
        self.fc2 = nn.Linear(hidden_dim, hidden_dim)
        # Non-linearity 2
        self.relu2 = nn.ReLU()

        # Linear function 3 (readout): 100 --> 10
        self.fc3 = nn.Linear(hidden_dim, output_dim)  

    def forward(self, x):
        # Linear function 1
        out = self.fc1(x)
        # Non-linearity 1
        out = self.relu1(out)

        # Linear function 2
        out = self.fc2(out)
        # Non-linearity 2
        out = self.relu2(out)

        # Linear function 3 (readout)
        out = self.fc3(out)
        return out
'''
STEP 4: INSTANTIATE MODEL CLASS
'''
input_dim = 28*28
hidden_dim = 100
output_dim = 10

model = FeedforwardNeuralNetModel(input_dim, hidden_dim, output_dim)

'''
STEP 5: INSTANTIATE LOSS CLASS
'''
criterion = nn.CrossEntropyLoss()


'''
STEP 6: INSTANTIATE OPTIMIZER CLASS
'''
learning_rate = 0.1

optimizer = torch.optim.SGD(model.parameters(), lr=learning_rate)

'''
STEP 7: TRAIN THE MODEL
'''
iter = 0
for epoch in range(num_epochs):
    for i, (images, labels) in enumerate(train_loader):
        # Load images with gradient accumulation capabilities
        images = images.view(-1, 28*28).requires_grad_()
        labels = labels

        # Clear gradients w.r.t. parameters
        optimizer.zero_grad()

        # Forward pass to get output/logits
        outputs = model(images)

        # Calculate Loss: softmax --> cross entropy loss
        loss = criterion(outputs, labels)

        # Getting gradients w.r.t. parameters
        loss.backward()

        # Updating parameters
        optimizer.step()

        iter += 1

        if iter % 500 == 0:
            # Calculate Accuracy         
            correct = 0
            total = 0
            # Iterate through test dataset
            for images, labels in test_loader:
                # Load images with gradient accumulation capabilities
                images = images.view(-1, 28*28).requires_grad_()

                # Forward pass only to get logits/output
                outputs = model(images)

                # Get predictions from the maximum value
                _, predicted = torch.max(outputs.data, 1)

                # Total number of labels
                total += labels.size(0)

                # Total correct predictions
                correct += (predicted == labels).sum()

            accuracy = 100 * correct / total

            # Print Loss
            print('Iteration: {}. Loss: {}. Accuracy: {}'.format(iter, loss.item(), accuracy))

Iteration: 500. Loss: 0.2943076193332672. Accuracy: 91
Iteration: 1000. Loss: 0.19378866255283356. Accuracy: 93
Iteration: 1500. Loss: 0.14716047048568726. Accuracy: 95
Iteration: 2000. Loss: 0.1172453910112381. Accuracy: 95
Iteration: 2500. Loss: 0.03787495568394661. Accuracy: 96
Iteration: 3000. Loss: 0.0711575448513031. Accuracy: 96


# Model E: 3 Hidden Layer Feedforward Neural Network(ReLU activation)

In [74]:
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets

"""
STEP 1: LOADING DATASET
"""
train_dataset = dsets.MNIST(root='./data',
                            train=True,
                            transform=transforms.ToTensor(),
                            download=True)

test_dataset = dsets.MNIST(root='./data',
                           train=False,
                           transform = transforms.ToTensor())

"""
STEP 2: MAKING DATASET ITERABLE
"""

batch_size = 100
n_iter = 3000
num_epochs = n_iters / (len(train_dataset) / batch_size)
num_epochs = int(num_epochs)

train_loader = torch.utils.data.DataLoader(dataset=train_dataset,
                                           batch_size=batch_size,
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset,
                                          batch_size=batch_size,
                                          shuffle=False)

"""
STEP 3: CREATE MODEL CLASS
"""

class FeedforwardNeuralNetModel(nn.Module):
  def __init__(self, input_dim, hidden_dim, output_dim):
    super(FeedforwardNeuralNetModel, self).__init__()
    # Linear function 1: 784 ---> 100
    self.fc1 = nn.Linear(input_dim, hidden_dim)
    # Non-linearity
    self.relu1 = nn.ReLU()

    # linear function 2: 100 ---> 100
    self.fc2 = nn.Linear(hidden_dim, hidden_dim)
    # Non-linearity
    self.relu2 = nn.ReLU()

    # Linear function 3: 100 ---> 100
    self.fc3 = nn.Linear(hidden_dim, hidden_dim)
    # Non-linearity
    self.relu3 = nn.ReLU()

    # Linear function 4 (readout): 100 ---> 10
    self.fc4 = nn.Linear(hidden_dim, output_dim)
    
  def forward(self, x):
    # linear function 1
    out = self.fc1(x)
    # Non-linearity 1
    out = self.relu1(out)

    # Linear function 2
    out = self.fc2(out)
    # Non-linearity 2
    out = self.relu2(out)

    # Linear function 3
    out = self.fc3(out)
    # Non-linearity 3 
    out = self.relu3(out)

    # Linear function 4 (readout)
    out = self.fc4(out)
    return out
  
"""
STEP 4: INSTANTIATE MODEL CLASS
"""

input_dim = 28*28
hidden_dim = 100
output_dim = 10

model = FeedforwardNeuralNetModel(input_dim, hidden_dim, output_dim)

"""
STEP 5: INSTANTIATE LOSS CLASS
"""

criterion = nn.CrossEntropyLoss()

"""
STEP 6: INSTANTIATE OPTIMIZER CLASS
"""
learning_rate = 0.1
optimizer = torch.optim.SGD(model.parameters(), lr=learning_rate)

"""
STEP 7: TRAIN THE MODEL
"""

iter = 0
for epoch in range(num_epochs):
  for i, (images, labels) in enumerate(train_loader):
    # load images with gradient accumulation capabilities
    images = images.view(-1, 28*28).requires_grad_()
    # clear gradients w.r.t. parameters
    optimizer.zero_grad()
    # Forward pass to get output/logits
    outputs = model(images)
    # Calculate Loss: softmax ---> cross entropy loss
    loss = criterion(outputs, labels)
    # Getting gradients w.r.t. parameters
    loss.backward()
    # Updating parameters
    optimizer.step()

    iter += 1
    if iter % 500 == 0:
      # Calculate Accuracy
      correct = 0
      total = 0
      # Iterate through test dataset
      for images, labels in test_loader:
        # load images with gradients accumulation capabilities
        images = images.view(-1, 28*28).requires_grad_()
        # Forward pass only to get logits/output
        outputs = model(images)
        # get predictions from maximum value
        _, predicted = torch.max(outputs.data, 1)
        # total number of label
        total += labels.size(0)
        # total correct predictions
        correct += (predicted == labels).sum()
      accuracy = 100 * correct / total
      print('Iteration : {}. Loss: {}. Accuracy: {}'.format(iter, loss.item(), accuracy))

Iteration : 500. Loss: 0.2598670721054077. Accuracy: 90
Iteration : 1000. Loss: 0.11560507863759995. Accuracy: 94
Iteration : 1500. Loss: 0.13716332614421844. Accuracy: 95
Iteration : 2000. Loss: 0.15583257377147675. Accuracy: 96
Iteration : 2500. Loss: 0.0744999572634697. Accuracy: 96
Iteration : 3000. Loss: 0.04102780669927597. Accuracy: 96
