In [7]:
import torch
import torch.nn.functional as F
from torch_geometric.datasets import Planetoid
from torch_geometric.nn import GATConv, GCNConv
import numpy as np

In [8]:
def load_data():
    dataset = Planetoid(root='/tmp/Cora', name='Cora')
    data = dataset[0]
    return data, dataset.num_features, dataset.num_classes


In [13]:
class GATNet(torch.nn.Module):
    def __init__(self, in_channels, hidden_channels, out_channels, heads=8):
        super(GATNet, self).__init__()
        self.conv1 = GATConv(in_channels, hidden_channels, heads=heads, dropout=0.6)
        self.conv2 = GATConv(hidden_channels * heads, out_channels, heads=1, concat=False, dropout=0.6)

    def forward(self, data):
        x, edge_index = data.x, data.edge_index
        x = F.dropout(x, p=0.6, training=self.training)
        x = self.conv1(x, edge_index)
        x = F.elu(x)
        x = F.dropout(x, p=0.6, training=self.training)
        x = self.conv2(x, edge_index)
        return F.log_softmax(x, dim=1)


In [14]:
def train(model, data, optimizer, criterion, epochs=200):
    model.train()
    for epoch in range(epochs):
        optimizer.zero_grad()
        out = model(data)
        loss = criterion(out[data.train_mask], data.y[data.train_mask])
        loss.backward()
        optimizer.step()
        
        if epoch % 10 == 0:
            val_acc = evaluate(model, data, data.val_mask)
            print(f'Epoch: {epoch:03d}, Loss: {loss:.4f}, Val Acc: {val_acc:.4f}')

In [15]:
def evaluate(model, data, mask):
    model.eval()
    with torch.no_grad():
        pred = model(data).argmax(dim=1)
        correct = (pred[mask] == data.y[mask]).sum()
        acc = int(correct) / int(mask.sum())
    return acc

In [16]:
def main():
    # Load data
    data, num_features, num_classes = load_data()
    device = torch.device('cuda' if torch.cuda.is_available() else 'cpu')
    data = data.to(device)
    
    # Initialize GAT model
    gat_model = GATNet(num_features, hidden_channels=8, out_channels=num_classes).to(device)
    optimizer = torch.optim.Adam(gat_model.parameters(), lr=0.005, weight_decay=5e-4)
    criterion = torch.nn.CrossEntropyLoss()
    
    # Train GAT
    print("Training GAT model...")
    train(gat_model, data, optimizer, criterion)
    
    # Evaluate GAT
    test_acc = evaluate(gat_model, data, data.test_mask)
    print(f'GAT Test Accuracy: {test_acc:.4f}')
    
    # Optional: Train and evaluate GCN
    print("\nTraining GCN model...")
    gcn_model = GCNNet(num_features, hidden_channels=16, out_channels=num_classes).to(device)
    optimizer = torch.optim.Adam(gcn_model.parameters(), lr=0.01, weight_decay=5e-4)
    
    train(gcn_model, data, optimizer, criterion)
    gcn_test_acc = evaluate(gcn_model, data, data.test_mask)
    print(f'GCN Test Accuracy: {gcn_test_acc:.4f}')
    
    # Comparison discussion
    print("\nComparison Discussion:")
    print("- GAT typically achieves slightly better accuracy than GCN due to its attention mechanism,")
    print("  which allows it to weigh neighbor contributions differently.")
    print("- GCN is computationally simpler and faster as it doesn't compute attention scores.")
    print("- GAT has more parameters due to multiple attention heads, increasing memory usage.")
    print("- The attention mechanism in GAT can provide better interpretability of node relationships.")

if __name__ == '__main__':
    main()

Training GAT model...
Epoch: 000, Loss: 2.0182, Val Acc: 0.3200
Epoch: 010, Loss: 0.4212, Val Acc: 0.7940
Epoch: 020, Loss: 0.0464, Val Acc: 0.7920
Epoch: 030, Loss: 0.0083, Val Acc: 0.7800
Epoch: 040, Loss: 0.0044, Val Acc: 0.7800
Epoch: 050, Loss: 0.0042, Val Acc: 0.7820
Epoch: 060, Loss: 0.0053, Val Acc: 0.7760
Epoch: 070, Loss: 0.0067, Val Acc: 0.7740
Epoch: 080, Loss: 0.0078, Val Acc: 0.7720
Epoch: 090, Loss: 0.0081, Val Acc: 0.7700
Epoch: 100, Loss: 0.0079, Val Acc: 0.7700
Epoch: 110, Loss: 0.0074, Val Acc: 0.7680
Epoch: 120, Loss: 0.0070, Val Acc: 0.7700
Epoch: 130, Loss: 0.0067, Val Acc: 0.7660
Epoch: 140, Loss: 0.0064, Val Acc: 0.7620
Epoch: 150, Loss: 0.0061, Val Acc: 0.7580
Epoch: 160, Loss: 0.0058, Val Acc: 0.7640
Epoch: 170, Loss: 0.0056, Val Acc: 0.7600
Epoch: 180, Loss: 0.0054, Val Acc: 0.7560
Epoch: 190, Loss: 0.0052, Val Acc: 0.7540
GAT Test Accuracy: 0.7960

Training GCN model...
Epoch: 000, Loss: 1.9552, Val Acc: 0.3540
Epoch: 010, Loss: 0.6521, Val Acc: 0.7760
Epoch