# Importing Libraries
performing all necessary imports here

In [1]:
# import model related libraries
from transformers import pipeline, AutoModelForCausalLM, AutoTokenizer

# import huggingface login
from huggingface_hub import login

# import the accesstoken of huggingface, which is saved in .env file
# from dotenv import load_dotenv # it was not running on kaggle so i have commented it
import os

# import time to compute how long it took the model to run
import time

# import date-time to display the date and time of last run
from datetime import datetime

# import text wrap to make sure its fully displayable
import textwrap

# Initialize Inputs
here we initailize all inputs given, like the input text, the prompt, and the model name

### Input Text
here is the email we will be passing for all models: this is given in the question and stays constant

In [2]:
input_text = """Subject: Concerns About Professor X’s Conduct 
 
Dear Dr. Ustaad, 
I hope this email finds you well. I am writing to express my concerns about Professor X’s conduct during the Introduction to Zoology class last semester. On multiple occasions, Professor X made dismissive remarks about students’ questions and failed to provide clear feedback on assignments.  
 
Additionally, the grading seemed inconsistent and unfair, with no opportunity for clarification or appeal. 
 
I found this experience deeply frustrating and demotivating, and I believe it affected my performance in the course. I would appreciate it if the department could look into this matter and ensure that future students have a more positive and supportive learning environment. 
 
Thank you for your attention to this matter. 
Sincerely, 
Shaagird
"""

print("Input Text:\n", input_text)

Input Text:
 Subject: Concerns About Professor X’s Conduct 
 
Dear Dr. Ustaad, 
I hope this email finds you well. I am writing to express my concerns about Professor X’s conduct during the Introduction to Zoology class last semester. On multiple occasions, Professor X made dismissive remarks about students’ questions and failed to provide clear feedback on assignments.  
 
Additionally, the grading seemed inconsistent and unfair, with no opportunity for clarification or appeal. 
 
I found this experience deeply frustrating and demotivating, and I believe it affected my performance in the course. I would appreciate it if the department could look into this matter and ensure that future students have a more positive and supportive learning environment. 
 
Thank you for your attention to this matter. 
Sincerely, 
Shaagird



### Model Params
here we intialize the model name, its parameters, and then the text_generator we will be using

In [3]:
model_name = "open-thoughts/OpenThinker-7B"
model = AutoModelForCausalLM.from_pretrained(
    model_name,
    device_map="auto",
    torch_dtype="auto",
    trust_remote_code=True,
)
tokenizer = AutoTokenizer.from_pretrained(model_name)

config.json:   0%|          | 0.00/733 [00:00<?, ?B/s]

model.safetensors.index.json:   0%|          | 0.00/27.8k [00:00<?, ?B/s]

Downloading shards:   0%|          | 0/4 [00:00<?, ?it/s]

model-00001-of-00004.safetensors:   0%|          | 0.00/4.88G [00:00<?, ?B/s]

model-00002-of-00004.safetensors:   0%|          | 0.00/4.93G [00:00<?, ?B/s]

model-00003-of-00004.safetensors:   0%|          | 0.00/4.33G [00:00<?, ?B/s]

model-00004-of-00004.safetensors:   0%|          | 0.00/1.09G [00:00<?, ?B/s]

Loading checkpoint shards:   0%|          | 0/4 [00:00<?, ?it/s]

generation_config.json:   0%|          | 0.00/243 [00:00<?, ?B/s]

tokenizer_config.json:   0%|          | 0.00/7.34k [00:00<?, ?B/s]

vocab.json:   0%|          | 0.00/2.78M [00:00<?, ?B/s]

merges.txt:   0%|          | 0.00/1.67M [00:00<?, ?B/s]

tokenizer.json:   0%|          | 0.00/11.4M [00:00<?, ?B/s]

added_tokens.json:   0%|          | 0.00/605 [00:00<?, ?B/s]

special_tokens_map.json:   0%|          | 0.00/616 [00:00<?, ?B/s]

In [4]:
generator = pipeline(
    "text-generation",
    model=model,
    tokenizer=tokenizer,
    return_full_text=False,
    max_new_tokens=5000,
    do_sample=False
)

Device set to use cuda:0


# Login, Authentication
here we authenticate the hugging face access token

In [5]:
# first we load the access token from the .env hidden file (it is a gitignore file)
# load_dotenv() # it was not running on kaggle so i have commentated it
# api_token = os.getenv("HUGGINGFACE_API_TOKEN")
api_token = "" # put your api token here, have removed for github security

In [6]:
login(api_token)

# Model Running
here we run the model, but before that we display the model details and also compute the timetaken to run the model

### Model Details
here we display how many parameters and how much memory the model took

In [7]:
print(model.dtype)

# Parameters Computation
total_params = sum(p.numel() for p in model.parameters())
print(f"Model Total Parameters: {total_params / 1e9:.2f} billion")

# Memory used Computation (in MBs)
memory = total_params * 2 / (1024 ** 2)  
print(f"Estimate Memory Footprint: {memory:.2f} MB\n")

torch.bfloat16
Model Total Parameters: 7.62 billion
Estimate Memory Footprint: 14525.64 MB



## Model Running
here for every task, we will first save the start_time and the end_time and then compute the timetaken

In [8]:
def compute_time(start_time, end_time):
    """This function calculates the time difference between two given times in seconds. For this given problem, this function computes the timetaken by a model to run and perform a task

    Args:
        start_time (time): start time of the model
        end_time (time): end time of the model

    Returns:
        string: the time taken in hours, minutes, seconds, and mili-seconds
    """
    elapsed_time = end_time - start_time
    hours, rem = divmod(elapsed_time, 3600)
    minutes, rem = divmod(rem, 60)
    seconds, milliseconds = divmod(rem, 1)
    milliseconds *= 1000  # Convert seconds fraction to milliseconds
    
    return f"Time taken: {int(hours)}h {int(minutes)}m {int(seconds)}s {int(milliseconds)}ms"

### Prompt Making
here we define the prompt we will use for all tasks of summarization, question answering, keyword extraction and translation

In [9]:
prompt = f"""
Perform the following actions: 
1 - Summarize the following text delimited by triple backticks.
2 - Answer “What was the reason for the student's disappointment with Professor X? ” based on the email content.
3 - Identify key details such as incidents, concerns, and requested actions. 
4 - Translate the following text into French.

Separate your answers with line breaks.
Text:
```{input_text}```
"""

### Model Execution
here we run the model by passing the prompt to a text_generator

In [10]:
# start the timer
start_time = time.time()

# generate the output
# messages = [
#     {"role": "user", "content": prompt}
# ]
output = generator(prompt)

# end the timer
end_time = time.time()

# display time taken with output
print(compute_time(start_time, end_time))
print("Output:\n", output[0]["generated_text"])#textwrap.fill(output[0]["generated_text"], width=80))



Time taken: 0h 1m 19s 304ms
Output:
 Answer 1:
The email expresses concerns about Professor X's conduct in the Introduction to Zoology class. The student describes dismissive remarks towards questions and inconsistent/unfair grading without opportunities for clarification or appeal. This led to frustration and decreased motivation, impacting their performance. They request an investigation to improve future student experiences.

Answer 2:
The student was disappointed with Professor X due to dismissive remarks about their questions and unfair, inconsistent grading without avenues for clarification or appeal, which demotivated them and harmed their course performance.

Answer 3:
Key details include:
- Multiple instances of dismissive remarks from Professor X.
- Inconsistent and unfair grading practices.
- Lack of opportunities for clarification or appeal regarding grades.
- Impact on the student's motivation and performance.
- Request for departmental intervention to address these issues

In [14]:
prompt_back_to_eng = """
Translate the following text to English:
Objet : Inquiétudes concernant le comportement du Professeur X  

Cher Dr. Ustaad,  

J'espère que ce courriel vous trouve en bonne santé. Je vous écris pour exprimer mes inquiétudes concernant le comportement du Professeur X lors de la classe d'introduction à la zoologie l'an dernier. À plusieurs reprises, le Professeur X a fait des remarques dédaigneuses envers les questions des étudiants et n'a pas fourni de retour sur leurs devoirs de manière claire et constructive.  

De plus, la notation semblait inconsistante et injuste, sans aucune opportunité de clarification ou d'appel.  

Cette expérience m'a troublé et a miné ma motivation, ce qui a affecté mon rendement dans le cours. J'apprécierais que le département examine cette situation et s'assure d'un environnement d'apprentissage plus positif et soutenant pour les futurs étudiants.  

Je vous remercie pour votre attention à ce sujet.  

Cordialement,  
Shaagird
"""

In [15]:
# start the timer
start_time = time.time()

# generate the output
# messages = [
#     {"role": "user", "content": prompt_back_to_eng}
# ]
output = generator(prompt_back_to_eng)

# end the timer
end_time = time.time()

# display time taken with output
print(compute_time(start_time, end_time))
print("Output:\n", output[0]["generated_text"])#textwrap.fill(output[0]["generated_text"], width=80))

Time taken: 0h 1m 19s 720ms
Output:
 Subject: Concerns about Professor X's Behavior

Dear Dr. Ustaad,

I hope this email finds you well. I am writing to express my concerns regarding the behavior of Professor X during last year's Introduction to Zoology class. On multiple occasions, Professor X made derogatory comments towards student questions and failed to provide clear and constructive feedback on their assignments.

Moreover, the grading appeared inconsistent and unfair, with no opportunity for clarification or appeal.

This experience has troubled me and diminished my motivation, impacting my performance in the course. I would appreciate it if the department could review this situation and ensure a more positive and supportive learning environment for future students.

Thank you for your attention to this matter.

Sincerely,
Student
<|begin_of_thought|>

Okay, let's see. The original text is in French, and I need to translate it into English. First, I should read through the entir

# End
here we display the last time this jupyter notebook was run to always remember it

In [16]:
print(datetime.now().strftime("%Y-%m-%d %H:%M:%S"))

2025-02-02 12:26:54
