In [1]:
import os
import json
import textwrap
import torch
import transformers

from dotenv import load_dotenv
from langchain.memory import ConversationBufferMemory
from langchain import LLMChain, PromptTemplate
from langchain import HuggingFacePipeline
from langchain import PromptTemplate,  LLMChain
from transformers import AutoTokenizer, AutoModelForCausalLM, pipeline
from huggingface_hub import login

In [2]:
load_dotenv()
login(token = os.getenv("HUGGINGFACE_TOKEN"))

True

## LLaMA2 7B Chat


In [4]:
tokenizer = AutoTokenizer.from_pretrained("meta-llama/Llama-2-7b-chat-hf",
                                          use_auth_token=True,)

model = AutoModelForCausalLM.from_pretrained("meta-llama/Llama-2-7b-chat-hf",
                                             device_map='auto',
                                             torch_dtype=torch.float16,
                                             use_auth_token=True,
                                            #  load_in_8bit=True,
                                            #  load_in_4bit=True
                                             )



Downloading (…)fetensors.index.json:   0%|          | 0.00/26.8k [00:00<?, ?B/s]

Downloading shards:   0%|          | 0/2 [00:00<?, ?it/s]

Downloading (…)of-00002.safetensors:   0%|          | 0.00/9.98G [00:00<?, ?B/s]

Downloading (…)of-00002.safetensors:   0%|          | 0.00/3.50G [00:00<?, ?B/s]

Loading checkpoint shards:   0%|          | 0/2 [00:00<?, ?it/s]



Downloading generation_config.json:   0%|          | 0.00/188 [00:00<?, ?B/s]



In [5]:
B_INST, E_INST = "[INST]", "[/INST]"
B_SYS, E_SYS = "<<SYS>>\n", "\n<</SYS>>\n\n"
DEFAULT_SYSTEM_PROMPT = """\
You are a helpful, respectful and honest assistant. Always answer as helpfully as possible, while being safe. Your answers should not include any harmful, unethical, racist, sexist, toxic, dangerous, or illegal content. Please ensure that your responses are socially unbiased and positive in nature.

If a question does not make any sense, or is not factually coherent, explain why instead of answering something not correct. If you don't know the answer to a question, please don't share false information."""



def get_prompt(instruction, new_system_prompt=DEFAULT_SYSTEM_PROMPT ):
    SYSTEM_PROMPT = B_SYS + new_system_prompt + E_SYS
    prompt_template =  B_INST + SYSTEM_PROMPT + instruction + E_INST
    return prompt_template

def cut_off_text(text, prompt):
    cutoff_phrase = prompt
    index = text.find(cutoff_phrase)
    if index != -1:
        return text[:index]
    else:
        return text

def remove_substring(string, substring):
    return string.replace(substring, "")



def generate(text):
    prompt = get_prompt(text)
    with torch.autocast('cuda', dtype=torch.bfloat16):
        inputs = tokenizer(prompt, return_tensors="pt").to('cuda')
        outputs = model.generate(**inputs,
                                 max_new_tokens=512,
                                 eos_token_id=tokenizer.eos_token_id,
                                 pad_token_id=tokenizer.eos_token_id,
                                 )
        final_outputs = tokenizer.batch_decode(outputs, skip_special_tokens=True)[0]
        final_outputs = cut_off_text(final_outputs, '</s>')
        final_outputs = remove_substring(final_outputs, prompt)

    return final_outputs#, outputs

def parse_text(text):
        wrapped_text = textwrap.fill(text, width=100)
        print(wrapped_text +'\n\n')
        # return assistant_text


In [6]:
pipe = pipeline("text-generation",
                model=model,
                tokenizer= tokenizer,
                torch_dtype=torch.bfloat16,
                device_map="auto",
                max_new_tokens = 512,
                do_sample=True,
                top_k=30,
                num_return_sequences=1,
                eos_token_id=tokenizer.eos_token_id
                )

In [7]:
instruction = "What is the temperature in Melbourne?"

get_prompt(instruction)

"[INST]<<SYS>>\nYou are a helpful, respectful and honest assistant. Always answer as helpfully as possible, while being safe. Your answers should not include any harmful, unethical, racist, sexist, toxic, dangerous, or illegal content. Please ensure that your responses are socially unbiased and positive in nature.\n\nIf a question does not make any sense, or is not factually coherent, explain why instead of answering something not correct. If you don't know the answer to a question, please don't share false information.\n<</SYS>>\n\nWhat is the temperature in Melbourne?[/INST]"

In [8]:
instruction = "Summarize the following text for me {text}"

system_prompt = "You are an expert and summarization and reducing the number of words used"

get_prompt(instruction, system_prompt)

'[INST]<<SYS>>\nYou are an expert and summarization and reducing the number of words used\n<</SYS>>\n\nSummarize the following text for me {text}[/INST]'

## LangChain basics

In [9]:
llm = HuggingFacePipeline(pipeline = pipe, model_kwargs = {'temperature':0})

In [10]:
system_prompt = "You are an advanced assistant that excels at translation. "
instruction = "Convert the following text from English to French:\n\n {text}"
template = get_prompt(instruction, system_prompt)
print(template)

prompt = PromptTemplate(template=template, input_variables=["text"])
llm_chain = LLMChain(prompt=prompt, llm=llm)

[INST]<<SYS>>
You are an advanced assistant that excels at translation. 
<</SYS>>

Convert the following text from English to French:

 {text}[/INST]


In [11]:
text = "how are you today?"
output = llm_chain.run(text)

parse_text(output)

  Bien sûr! Here is the translation of "how are you today?" in French:  Comment allez-vous
aujourd'hui?




### Summarization

In [12]:
instruction = "Summarize the following article for me {text}"
system_prompt = "You are an expert in summarization and expressing key ideas succintly"

template = get_prompt(instruction, system_prompt)
print(template)

prompt = PromptTemplate(template=template, input_variables=["text"])
llm_chain = LLMChain(prompt=prompt, llm=llm)

[INST]<<SYS>>
You are an expert in summarization and expressing key ideas succintly
<</SYS>>

Summarize the following article for me {text}[/INST]


In [13]:
def count_words(input_string):
    words = input_string.split(" ")
    return len(words)

text = '''Twitter (now X) CEO Linda Yaccarino claims usage at ‘all time high’ in memo to staff
Twitter’s (now X’s) newly established CEO Linda Yaccarino touts the company’s success and X’s future plans in a company-wide memo obtained by CNBC. The exec once again claims, without sharing any specific metrics, that the service’s usage is at an “all time high,” and hints at what’s to come in terms of new product experiences for the newly rebranded platform.

The service formerly known as Twitter has been working to become more than just a social network and more of an “everything app,” as owner Elon Musk dubbed it.

As the Telsa and Space X exec explained in October 2022, telegraphing Twitter’s eventual rebranding, buying Twitter was meant to be “an accelerant to creating X, the everything app.”


His grand plan has been to create an app that allows creators to monetize their content, then later moves into payments services and even banking, Musk remarked during a Twitter Spaces livestream with advertisers in November. At the time, he even mentioned the possibility of establishing money market accounts on Twitter that would pay a high-interest rate to attract consumers to X.

Those possible product concepts were again referenced in Yaccarino’s new missive, when she writes, “Our usage is at an all time high and we’ll continue to delight our entire community with new experiences in audio, video, messaging, payments, banking – creating a global marketplace for ideas, goods, services, and opportunities.”

Twitter, now X, has already implemented some of Musk’s ideas around videos and creator monetization. In May, the company began allowing subscribers to upload two-hour videos to its service, which advertiser Apple then leveraged when it released the entire first episode of its hit Apple TV+ show “Silo” on the platform. Fired Fox News host Tucker Carlson had been posting lengthy videos to Twitter as well, until ordered to stop by the network.

In addition, earlier this month, Twitter began sharing ad revenue with verified creators.

However, all is not well at Twitter X, whose traffic — at least by third-party measurements — has been dropping. Data from web analytics firm Similarweb indicated Twitter’s web traffic declined 5% for the first two days its latest rival, Instagram Threads, became generally available, compared with the week prior. Plus, Similarweb said Twitter’s web traffic was down 11% compared with the same days in 2022. Additionally, Cloudflare CEO Matthew Prince earlier this month tweeted a graph of traffic to the Twitter.com domain that showed “Twitter traffic tanking,” he said.


Yaccarino subtly pushed back at those reports at the time, claiming that Twitter had its largest usage day since February in early July. She did not share any specific metrics or data. At the same time, however, the company was quietly blocking links to Threads.net in Twitter searches, suggesting it was concerned about the new competition.

Today, Yaccarino repeats her vague claims around X’s high usage in her company-wide memo even as analysts at Forrester are predicting X will either shut down or be acquired within the next 12 months and numerous critics concur that the X rebrand is destined to fail.

Yaccarino’s memo, otherwise, was mostly a lot of cheerleading, applauding X’s team for their work and touting X’s ability to “impress the world all over again,” as Twitter once did.

The full memo, courtesy of CBNC, is below:

Hi team,

What a momentous weekend. As I said yesterday, it’s extremely rare, whether it’s in life or in business, that you have the opportunity to make another big impression. That’s what we’re experiencing together, in real time. Take a moment to put it all into perspective.

17 years ago, Twitter made a lasting imprint on the world. The platform changed the speed at which people accessed information. It created a new dynamic for how people communicated, debated, and responded to things happening in the world. Twitter introduced a new way for people, public figures, and brands to build long lasting relationships. In one way or another, everyone here is a driving force in that change. But equally all our users and partners constantly challenged us to dream bigger, to innovate faster, and to fulfill our great potential.

With X we will go even further to transform the global town square — and impress the world all over again.

Our company uniquely has the drive to make this possible. Many companies say they want to move fast — but we enjoy moving at the speed of light, and when we do, that’s X. At our core, we have an inventor mindset — constantly learning, testing out new approaches, changing to get it right and ultimately succeeding.

With X, we serve our entire community of users and customers by working tirelessly to preserve free expression and choice, create limitless interactivity, and create a marketplace that enables the economic success of all its participants.

The best news is we’re well underway. Everyone should be proud of the pace of innovation over the last nine months — from long form content, to creator monetization, and tremendous advancements in brand safety protections. Our usage is at an all time high and we’ll continue to delight our entire community with new experiences in audio, video, messaging, payments, banking – creating a global marketplace for ideas, goods, services, and opportunities.

Please don’t take this moment for granted. You’re writing history, and there’s no limit to our transformation. And everyone, is invited to build X with us.

Elon and I will be working across every team and partner to bring X to the world. That includes keeping our entire community up to date, ensuring that we all have the information we need to move forward.

Now, let’s go make that next big impression on the world, together.

Linda'''

count_words(text)

940

In [14]:
output = llm_chain.run(text)
print(count_words(output))
parse_text(output)

98
  Twitter CEO Linda Yaccarino has claimed in a company-wide memo that the platform's usage is at an
"all time high," despite recent traffic declines. Yaccarino touted the company's success and future
plans, including the creation of a "global marketplace for ideas, goods, services, and
opportunities" through new product experiences. The memo comes as analysts predict X will either
shut down or be acquired within the next 12 months, and numerous critics concur that the X rebrand
is destined to fail. Yaccarino's memo largely consisted of cheerleading and encouraging the team to
keep innovating and making history.




## Simple Chatbot

In [15]:
instruction = "Chat History:\n\n{chat_history} \n\nUser: {user_input}"
system_prompt = "You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context"

template = get_prompt(instruction, system_prompt)
print(template)

[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

{chat_history} 

User: {user_input}[/INST]


In [16]:
prompt = PromptTemplate(
    input_variables=["chat_history", "user_input"], template=template
)
memory = ConversationBufferMemory(memory_key="chat_history")

In [17]:
llm_chain = LLMChain(
    llm=llm,
    prompt=prompt,
    verbose=True,
    memory=memory,
)

In [18]:
llm_chain.predict(user_input="Hi, my name is Sam")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

 

User: Hi, my name is Sam[/INST][0m

[1m> Finished chain.[0m


"  Hello Sam! I'm here to help. How can I assist you today?"

In [19]:
llm_chain.predict(user_input="Can you tell me about yourself.")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today? 

User: Can you tell me about yourself.[/INST][0m

[1m> Finished chain.[0m


"  Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss?"

In [20]:
llm_chain.predict(user_input="Today is Friday. What number day of the week is that?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss? 

User: Today is Friday. What number day of the week is that?[/INST][0m

[1m> Finished chain.[0m


'  AI: Great question! Friday is the 5th day of the week.'

In [21]:
llm_chain.predict(user_input="what is the day today?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week. 

User: what is the day today?[/INST][0m

[1m> Finished chain.[0m


'  AI:  Great question! *smiling* Today is Friday, the 5th day of the week.'

In [22]:
llm_chain.predict(user_input="What is my name?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI:  Great question! *smiling* Today is Friday, the 5th day of the week. 

User: What is my name?[/INST][0m

[1m> Finished chain.[0m


'  AI: Hello! *smiling* Your name is Sam.'

In [23]:
llm_chain.predict(user_input="Can you tell me about the olympics")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI:  Great question! *smiling* Today is Friday, the 5th day of the week.
Human: What is my name?
AI:   AI: Hello! *smiling* Your name is Sam. 

User: Can you tell me about the olympics[/INST][0m

[1m> Finished chain.[0m


'  Of course! The Olympics are a major international sporting event held every four years, where athletes from around the world compete in various sports such as track and field, swimming, gymnastics, and many more. The Olympics have a rich history dating back to 1896, and they have evolved over the years to include new sports and events. The Olympics are also a symbol of international unity and cooperation, bringing together athletes and spectators from all over the world. Is there anything else you would like to know about the Olympics?'

In [24]:
llm_chain.predict(user_input="What have we talked about in this Chat?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! I'm here to help. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! *smiling* I'm just an AI, I don't have a personal life or experiences like humans do, but I'm here to help you with any questions or tasks you may have. Is there something specific you'd like to know or discuss?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI:  Great question! *smiling* Today is Friday, the 5th day of the week.
Human: What is my name?
AI:   AI: Hello! *smiling* Your name is Sam.
Human: Can you tell me about the olympics
AI:   Of course! The Olympics are a major international spo

"  Sure, I can help you with that! Based on the chat history you provided, we have talked about:\n\n1. The assistant's name and greeting.\n2. The user's name (Sam).\n3. The day of the week (Friday).\n4. The user's question about the Olympics.\n\nIs there anything else you would like to know or discuss?"