In [23]:
from numpy import *

In [24]:
def compute_error_for_line_given_points(b,m,points):
    totalError = 0 # initialize at 0
    # for every point
    for i in range(0, len(points)):
        # get x value
        x = points[i, 0]
        # get y value
        y = points[i, 1]
        # get the difference, square it and add it to the total
        totalError += (y - (m * x + b)) ** 2
        
    return totalError / float(len(points))

In [25]:
def step_gradient(b_current, m_current, points, learningRate):
    # these are the starting points for our gradient (gradient => slope)
    b_gradient = 0
    m_gradient = 0
    
    N = float(len(points))
    for i in range(0, len(points)):
        x = points[i, 0]
        y = points[i, 1]
        # direction with respect to b and m
        # computing partial derivates for our error function
        b_gradient += -(2/N) * (y - ((m_current * x) + b_current))
        m_gradient += (2/N) * x * (y - ((m_current * x) + b_current))
    
    # update our b and m values using our partial derivatives
    new_b = b_current - (learningRate * b_gradient)
    new_m = m_current - (learningRate * m_gradient)
    
    return [new_b, new_m]

In [26]:
def gradient_descent_runner(points, starting_b, starting_m, learning_rate, num_iterations):
    # starting b and m
    b = starting_b
    m = starting_m
    
    # gradient descent 
    for i in range(num_iterations):
        # update b and m with the new more accurate b and m by performing this gradient step
        b, m = step_gradient(b, m, array(points), learning_rate)
    return [b, m]

In [27]:
def run():
    points = genfromtxt("data.csv", delimiter=",")
    learning_rate = 0.0001
    initial_b = 0 # initial y-intercept guess
    initial_m = 0 # initial slope guess
    num_iterations = 1000
    print("Starting gradient descent at b = {0}, m = {1}, error = {2}".format(initial_b, initial_m, compute_error_for_line_given_points(initial_b, initial_m, points)))
    print("Running...")
    [b, m] = gradient_descent_runner(points, initial_b, initial_m, learning_rate, num_iterations)
    print("After {0} iterations b = {1}, m = {2}, error = {3}".format(num_iterations, b, m, compute_error_for_line_given_points(b, m, points)))
    

In [28]:
run()

Starting gradient descent at b = 0, m = 0, error = 5565.107834483211
Running...
After 1000 iterations b = 9.451159314734957e+173, m = -4.808747047443573e+175, error = inf


