<a href="https://colab.research.google.com/github/tharun-0-0-6/sdc/blob/main/using_data_sets.ipynb" target="_parent"><img src="https://colab.research.google.com/assets/colab-badge.svg" alt="Open In Colab"/></a>

In [None]:
#Logistic Regression on Titanic Dataset
# 📦 Step 1: Import necessary libraries
import numpy as np
import pandas as pd
import seaborn as sns
import matplotlib.pyplot as plt
from sklearn.model_selection import train_test_split
from sklearn.preprocessing import StandardScaler
from sklearn.linear_model import LogisticRegression
from sklearn.metrics import classification_report, confusion_matrix, accuracy_score

# 🧠 Step 2: Load Titanic dataset
url = "https://raw.githubusercontent.com/datasciencedojo/datasets/master/titanic.csv"
df = pd.read_csv(url)

# 👀 Step 3: View sample data
print("Sample Data:\n", df.head())

# 🔍 Step 4: Preprocess data
df = df[["Survived", "Pclass", "Sex", "Age", "Fare"]].dropna()  # Select relevant features and drop missing values

# Convert categorical variable 'Sex' to numeric (0 = male, 1 = female)
df["Sex"] = df["Sex"].map({"male": 0, "female": 1})

# ✅ Step 5: Split into features (X) and target (y)
X = df.drop("Survived", axis=1)
y = df["Survived"]

# 🧪 Step 6: Train-test split
X_train, X_test, y_train, y_test = train_test_split(X, y, test_size=0.2, random_state=42)

# 🔢 Step 7: Standardize the data
scaler = StandardScaler()
X_train = scaler.fit_transform(X_train)
X_test = scaler.transform(X_test)

# 🔢 Step 8: Train Logistic Regression model
model = LogisticRegression(max_iter=1000)
model.fit(X_train, y_train)

# 🔮 Step 9: Predict
y_pred = model.predict(X_test)

# 📊 Step 10: Evaluate model
print("\n📈 Classification Report:")
print(classification_report(y_test, y_pred))

print("\n📊 Confusion Matrix:")
print(confusion_matrix(y_test, y_pred))

print("\n✅ Accuracy Score:", accuracy_score(y_test, y_pred))

# 🧪 Step 11: Real-time prediction example
print("\n🔍 Real-Time Prediction Example")
pclass = int(input("Enter passenger class (1, 2, or 3): "))
sex = int(input("Enter sex (0 for male, 1 for female): "))
age = float(input("Enter age: "))
fare = float(input("Enter ticket fare: "))

# Scale input
sample_input = scaler.transform([[pclass, sex, age, fare]])

# Predict survival
predicted = model.predict(sample_input)[0]
label = "Survived" if predicted == 1 else "Did Not Survive"
print(f"🛳 Prediction: {label}")


Sample Data:
    PassengerId  Survived  Pclass  \
0            1         0       3   
1            2         1       1   
2            3         1       3   
3            4         1       1   
4            5         0       3   

                                                Name     Sex   Age  SibSp  \
0                            Braund, Mr. Owen Harris    male  22.0      1   
1  Cumings, Mrs. John Bradley (Florence Briggs Th...  female  38.0      1   
2                             Heikkinen, Miss. Laina  female  26.0      0   
3       Futrelle, Mrs. Jacques Heath (Lily May Peel)  female  35.0      1   
4                           Allen, Mr. William Henry    male  35.0      0   

   Parch            Ticket     Fare Cabin Embarked  
0      0         A/5 21171   7.2500   NaN        S  
1      0          PC 17599  71.2833   C85        C  
2      0  STON/O2. 3101282   7.9250   NaN        S  
3      0            113803  53.1000  C123        S  
4      0            373450   8.0500   NaN

