# 1. Importing all the required modules

In [1]:
import numpy as np
import matplotlib.pyplot as plt
%matplotlib inline
import torch
import torch.nn as nn
import torchvision.transforms as transforms
import torchvision.datasets as dsets
from torch.autograd import Variable
import collections

# 2. Load the dataset

In [2]:
train_dataset = dsets.MNIST(root='./data',
                            train=True,
                            transform=transforms.ToTensor(),
                            download=True)

test_dataset = dsets.MNIST(root='./data',
                           train=False,
                           transform=transforms.ToTensor())

# 3. Set Parameters 

In [3]:
batch_size = 100
n_iters = 3000
num_epochs = int(n_iters / (len(train_dataset) / batch_size))

# 4. Make the dataset iterable 

In [4]:
train_loader = torch.utils.data.DataLoader(dataset=train_dataset,
                                           batch_size=batch_size,
                                           shuffle=True)

test_loader = torch.utils.data.DataLoader(dataset=test_dataset,
                                           batch_size=batch_size,
                                           shuffle=False)

# 5. Define Convolutional Neural Network Model

In [5]:
class CNNModel(nn.Module):
    def __init__(self):
        super(CNNModel,self).__init__()
        
        self.cnn1 = nn.Conv2d(in_channels=1,out_channels=16,kernel_size=5,stride=1,padding=2)
        self.relu1 = nn.ReLU()
        self.avgpool1 = nn.AvgPool2d(kernel_size=2)
        
        self.cnn2 = nn.Conv2d(in_channels=16,out_channels=32,kernel_size=5,stride=1,padding=2)
        self.relu2 = nn.ReLU()
        self.avgpool2 = nn.AvgPool2d(kernel_size=2)
        
        self.fc1 = nn.Linear(32*7*7,10)
        
    def forward(self,x):
        out = self.cnn1(x)
        out = self.relu1(out)
        out = self.avgpool1(out)
        out = self.cnn2(out)
        out = self.relu2(out)
        out = self.avgpool2(out)
        out = out.view(out.size(0),-1)
        out = self.fc1(out)
        
        return out

# 6. Instantiate model, criterion and optimizer

In [6]:
learning_rate = 0.1

model = CNNModel()
if torch.cuda.is_available():
    model.cuda()
criterion = nn.CrossEntropyLoss()
optimizer = torch.optim.SGD(model.parameters(),lr=learning_rate)

#  7. Train the model

In [7]:
iter = 0
for epoch in range(num_epochs):
    for i, (images,labels) in enumerate(train_loader):
        if torch.cuda.is_available():
            images = Variable(images.cuda())
            labels = Variable(labels.cuda())
        else:
            images = Variable(images)
            labels = Variable(labels)
        
        optimizer.zero_grad()
        outputs = model(images)
        loss = criterion(outputs, labels)
        loss.backward()
        optimizer.step()
        
        iter += 1
        if iter%500 == 0:
            correct = 0
            total = 0
            
            for images, labels in test_loader:
                if torch.cuda.is_available():
                    images = Variable(images.cuda())
                else:
                    images = Variable(images)
                outputs = model(images)
                
                _,predicted = torch.max(outputs.data,1)
                total += labels.size(0)
                correct += (predicted.cpu() == labels.cpu()).sum()
            
            accuracy = 100 * correct / total
            print('Iteration:{}. Loss:{}. Accuracy:{}'.format(iter,loss.data[0],accuracy))



Iteration:500. Loss:0.14340238273143768. Accuracy:95
Iteration:1000. Loss:0.0858822613954544. Accuracy:97
Iteration:1500. Loss:0.050326038151979446. Accuracy:97
Iteration:2000. Loss:0.14874526858329773. Accuracy:97
Iteration:2500. Loss:0.05480022728443146. Accuracy:97
Iteration:3000. Loss:0.02227490395307541. Accuracy:98


# Save trained model

In [8]:
torch.save(model.state_dict(),'CNN-model.pkl')