In [1]:
!pip -q install git+https://github.com/huggingface/transformers # need to install from github
!pip install -q datasets loralib sentencepiece
!pip -q install bitsandbytes accelerate xformers einops
!pip -q install langchain

In [2]:
!nvidia-smi

Fri Oct 20 14:39:10 2023       
+-----------------------------------------------------------------------------+
| NVIDIA-SMI 525.60.13    Driver Version: 525.60.13    CUDA Version: 12.0     |
|-------------------------------+----------------------+----------------------+
| GPU  Name        Persistence-M| Bus-Id        Disp.A | Volatile Uncorr. ECC |
| Fan  Temp  Perf  Pwr:Usage/Cap|         Memory-Usage | GPU-Util  Compute M. |
|                               |                      |               MIG M. |
|   0  NVIDIA TITAN RTX    On   | 00000000:08:00.0 Off |                  N/A |
| 41%   33C    P8    14W / 280W |     15MiB / 24576MiB |      0%      Default |
|                               |                      |                  N/A |
+-------------------------------+----------------------+----------------------+
                                                                               
+-----------------------------------------------------------------------------+
| Proces

In [3]:
from huggingface_hub import notebook_login

notebook_login()


VBox(children=(HTML(value='<center> <img\nsrc=https://huggingface.co/front/assets/huggingface_logo-noborder.sv‚Ä¶

## LLaMA2 7B Chat


In [5]:
import torch
import transformers
from transformers import AutoTokenizer, AutoModelForCausalLM

In [6]:

tokenizer = AutoTokenizer.from_pretrained("meta-llama/Llama-2-7b-chat-hf",
                                          use_auth_token=True,)

model = AutoModelForCausalLM.from_pretrained("meta-llama/Llama-2-7b-chat-hf",
                                             device_map='auto',
                                             torch_dtype=torch.float16,
                                             use_auth_token=True,
                                            #  load_in_8bit=True,
                                            #  load_in_4bit=True
                                             )



Loading checkpoint shards:   0%|          | 0/2 [00:00<?, ?it/s]



In [7]:
# Use a pipeline for later
from transformers import pipeline

pipe = pipeline("text-generation",
                model=model,
                tokenizer= tokenizer,
                torch_dtype=torch.bfloat16,
                device_map="auto",
                max_new_tokens = 512,
                do_sample=True,
                top_k=30,
                num_return_sequences=1,
                eos_token_id=tokenizer.eos_token_id
                )

In [8]:
!nvidia-smi

Fri Oct  6 14:03:56 2023       
+-----------------------------------------------------------------------------+
| NVIDIA-SMI 525.60.13    Driver Version: 525.60.13    CUDA Version: 12.0     |
|-------------------------------+----------------------+----------------------+
| GPU  Name        Persistence-M| Bus-Id        Disp.A | Volatile Uncorr. ECC |
| Fan  Temp  Perf  Pwr:Usage/Cap|         Memory-Usage | GPU-Util  Compute M. |
|                               |                      |               MIG M. |
|   0  NVIDIA TITAN RTX    On   | 00000000:08:00.0 Off |                  N/A |
| 41%   37C    P2    62W / 280W |  13482MiB / 24576MiB |      0%      Default |
|                               |                      |                  N/A |
+-------------------------------+----------------------+----------------------+
                                                                               
+-----------------------------------------------------------------------------+
| Proces

In [9]:
import json
import textwrap

B_INST, E_INST = "[INST]", "[/INST]"
B_SYS, E_SYS = "<<SYS>>\n", "\n<</SYS>>\n\n"
DEFAULT_SYSTEM_PROMPT = """\
You are a helpful, respectful and honest assistant. Always answer as helpfully as possible, while being safe. Your answers should not include any harmful, unethical, racist, sexist, toxic, dangerous, or illegal content. Please ensure that your responses are socially unbiased and positive in nature.

If a question does not make any sense, or is not factually coherent, explain why instead of answering something not correct. If you don't know the answer to a question, please don't share false information."""



def get_prompt(instruction, new_system_prompt=DEFAULT_SYSTEM_PROMPT ):
    SYSTEM_PROMPT = B_SYS + new_system_prompt + E_SYS
    prompt_template =  B_INST + SYSTEM_PROMPT + instruction + E_INST
    return prompt_template

def cut_off_text(text, prompt):
    cutoff_phrase = prompt
    index = text.find(cutoff_phrase)
    if index != -1:
        return text[:index]
    else:
        return text

def remove_substring(string, substring):
    return string.replace(substring, "")



def generate(text):
    prompt = get_prompt(text)
    with torch.autocast('cuda', dtype=torch.bfloat16):
        inputs = tokenizer(prompt, return_tensors="pt").to('cuda')
        outputs = model.generate(**inputs,
                                 max_new_tokens=512,
                                 eos_token_id=tokenizer.eos_token_id,
                                 pad_token_id=tokenizer.eos_token_id,
                                 )
        final_outputs = tokenizer.batch_decode(outputs, skip_special_tokens=True)[0]
        final_outputs = cut_off_text(final_outputs, '</s>')
        final_outputs = remove_substring(final_outputs, prompt)

    return final_outputs#, outputs

def parse_text(text):
        wrapped_text = textwrap.fill(text, width=100)
        print(wrapped_text +'\n\n')
        # return assistant_text


In [10]:
instruction = "What is the temperature in Melbourne?"

get_prompt(instruction)

"[INST]<<SYS>>\nYou are a helpful, respectful and honest assistant. Always answer as helpfully as possible, while being safe. Your answers should not include any harmful, unethical, racist, sexist, toxic, dangerous, or illegal content. Please ensure that your responses are socially unbiased and positive in nature.\n\nIf a question does not make any sense, or is not factually coherent, explain why instead of answering something not correct. If you don't know the answer to a question, please don't share false information.\n<</SYS>>\n\nWhat is the temperature in Melbourne?[/INST]"

In [11]:
instruction = "Summarize the following text for me {text}"

system_prompt = "You are an expert and summarization and reducing the number of words used"

get_prompt(instruction, system_prompt)

'[INST]<<SYS>>\nYou are an expert and summarization and reducing the number of words used\n<</SYS>>\n\nSummarize the following text for me {text}[/INST]'

## LangChain basics

In [12]:
from langchain import HuggingFacePipeline
from langchain import PromptTemplate,  LLMChain


In [13]:
llm = HuggingFacePipeline(pipeline = pipe, model_kwargs = {'temperature':0})

In [14]:
system_prompt = "You are an advanced assistant that excels at translation. "
instruction = "Convert the following text from English to French:\n\n {text}"
template = get_prompt(instruction, system_prompt)
print(template)

prompt = PromptTemplate(template=template, input_variables=["text"])
llm_chain = LLMChain(prompt=prompt, llm=llm)

[INST]<<SYS>>
You are an advanced assistant that excels at translation. 
<</SYS>>

Convert the following text from English to French:

 {text}[/INST]


In [15]:
text = "how are you today?"
output = llm_chain.run(text)

parse_text(output)

  Bien s√ªr! Here is the translation of "how are you today?" from English to French:  Comment √ßa va
aujourd'hui?




### Summarization

In [None]:
instruction = "Summarize the following article for me {text}"
system_prompt = "You are an expert and summarization and expressing key ideas succintly"

template = get_prompt(instruction, system_prompt)
print(template)

prompt = PromptTemplate(template=template, input_variables=["text"])
llm_chain = LLMChain(prompt=prompt, llm=llm)

[INST]<<SYS>>
You are an expert and summarization and expressing key ideas succintly
<</SYS>>

Summarize the following article for me {text}[/INST]


In [None]:
def count_words(input_string):
    words = input_string.split(" ")
    return len(words)

text = '''Twitter (now X) CEO Linda Yaccarino claims usage at ‚Äòall time high‚Äô in memo to staff
Twitter‚Äôs (now X‚Äôs) newly established CEO Linda Yaccarino touts the company‚Äôs success and X‚Äôs future plans in a company-wide memo obtained by CNBC. The exec once again claims, without sharing any specific metrics, that the service‚Äôs usage is at an ‚Äúall time high,‚Äù and hints at what‚Äôs to come in terms of new product experiences for the newly rebranded platform.

The service formerly known as Twitter has been working to become more than just a social network and more of an ‚Äúeverything app,‚Äù as owner Elon Musk dubbed it.

As the Telsa and Space X exec explained in October 2022, telegraphing Twitter‚Äôs eventual rebranding, buying Twitter was meant to be ‚Äúan accelerant to creating X, the everything app.‚Äù


His grand plan has been to create an app that allows creators to monetize their content, then later moves into payments services and even banking, Musk remarked during a Twitter Spaces livestream with advertisers in November. At the time, he even mentioned the possibility of establishing money market accounts on Twitter that would pay a high-interest rate to attract consumers to X.

Those possible product concepts were again referenced in Yaccarino‚Äôs new missive, when she writes, ‚ÄúOur usage is at an all time high and we‚Äôll continue to delight our entire community with new experiences in audio, video, messaging, payments, banking ‚Äì creating a global marketplace for ideas, goods, services, and opportunities.‚Äù

Twitter, now X, has already implemented some of Musk‚Äôs ideas around videos and creator monetization. In May, the company began allowing subscribers to upload two-hour videos to its service, which advertiser Apple then leveraged when it released the entire first episode of its hit Apple TV+ show ‚ÄúSilo‚Äù on the platform. Fired Fox News host Tucker Carlson had been posting lengthy videos to Twitter as well, until ordered to stop by the network.

In addition, earlier this month, Twitter began sharing ad revenue with verified creators.

However, all is not well at Twitter X, whose traffic ‚Äî at least by third-party measurements ‚Äî has been dropping. Data from web analytics firm Similarweb indicated Twitter‚Äôs web traffic declined 5% for the first two days its latest rival, Instagram Threads, became generally available, compared with the week prior. Plus, Similarweb said Twitter‚Äôs web traffic was down 11% compared with the same days in 2022. Additionally, Cloudflare CEO Matthew Prince earlier this month tweeted a graph of traffic to the Twitter.com domain that showed ‚ÄúTwitter traffic tanking,‚Äù he said.


Yaccarino subtly pushed back at those reports at the time, claiming that Twitter had its largest usage day since February in early July. She did not share any specific metrics or data. At the same time, however, the company was quietly blocking links to Threads.net in Twitter searches, suggesting it was concerned about the new competition.

Today, Yaccarino repeats her vague claims around X‚Äôs high usage in her company-wide memo even as analysts at Forrester are predicting X will either shut down or be acquired within the next 12 months and numerous critics concur that the X rebrand is destined to fail.

Yaccarino‚Äôs memo, otherwise, was mostly a lot of cheerleading, applauding X‚Äôs team for their work and touting X‚Äôs ability to ‚Äúimpress the world all over again,‚Äù as Twitter once did.

The full memo, courtesy of CBNC, is below:

Hi team,

What a momentous weekend. As I said yesterday, it‚Äôs extremely rare, whether it‚Äôs in life or in business, that you have the opportunity to make another big impression. That‚Äôs what we‚Äôre experiencing together, in real time. Take a moment to put it all into perspective.

17 years ago, Twitter made a lasting imprint on the world. The platform changed the speed at which people accessed information. It created a new dynamic for how people communicated, debated, and responded to things happening in the world. Twitter introduced a new way for people, public figures, and brands to build long lasting relationships. In one way or another, everyone here is a driving force in that change. But equally all our users and partners constantly challenged us to dream bigger, to innovate faster, and to fulfill our great potential.

With X we will go even further to transform the global town square ‚Äî and impress the world all over again.

Our company uniquely has the drive to make this possible. Many companies say they want to move fast ‚Äî but we enjoy moving at the speed of light, and when we do, that‚Äôs X. At our core, we have an inventor mindset ‚Äî constantly learning, testing out new approaches, changing to get it right and ultimately succeeding.

With X, we serve our entire community of users and customers by working tirelessly to preserve free expression and choice, create limitless interactivity, and create a marketplace that enables the economic success of all its participants.

The best news is we‚Äôre well underway. Everyone should be proud of the pace of innovation over the last nine months ‚Äî from long form content, to creator monetization, and tremendous advancements in brand safety protections. Our usage is at an all time high and we‚Äôll continue to delight our entire community with new experiences in audio, video, messaging, payments, banking ‚Äì creating a global marketplace for ideas, goods, services, and opportunities.

Please don‚Äôt take this moment for granted. You‚Äôre writing history, and there‚Äôs no limit to our transformation. And everyone, is invited to build X with us.

Elon and I will be working across every team and partner to bring X to the world. That includes keeping our entire community up to date, ensuring that we all have the information we need to move forward.

Now, let‚Äôs go make that next big impression on the world, together.

Linda'''

count_words(text)

940

In [None]:
output = llm_chain.run(text)
print(count_words(output))
parse_text(output)

124
  Sure, here is a summary of the article in 400 words or less: Twitter's newly appointed CEO, Linda
Yaccarino, has claimed in a memo to staff that the platform's usage is at an all-time high. She also
highlighted the company's grand plans to expand beyond social media and become an "everything app"
as envisioned by owner Elon Musk. Yaccarino touted the company's innovation in areas such as long-
form content, creator monetization, and brand safety protections. Despite reports of traffic
decline, Yaccarino remained optimistic and emphasized the company's commitment to preserving free
expression and creating a global marketplace for ideas and opportunities. She also acknowledged the
challenges ahead and encouraged the team to work together to make the next big impression on the
world.




## Simple Chatbot

In [None]:
from langchain.memory import ConversationBufferMemory
from langchain import LLMChain, PromptTemplate

In [None]:
instruction = "Chat History:\n\n{chat_history} \n\nUser: {user_input}"
system_prompt = "You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context"

template = get_prompt(instruction, system_prompt)
print(template)

[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

{chat_history} 

User: {user_input}[/INST]


In [None]:

prompt = PromptTemplate(
    input_variables=["chat_history", "user_input"], template=template
)
memory = ConversationBufferMemory(memory_key="chat_history")

In [None]:
llm_chain = LLMChain(
    llm=llm,
    prompt=prompt,
    verbose=True,
    memory=memory,
)

In [None]:
llm_chain.predict(user_input="Hi, my name is Sam")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

 

User: Hi, my name is Sam[/INST][0m

[1m> Finished chain.[0m


"  Hello Sam! It's nice to meet you. How can I assist you today?"

In [None]:
llm_chain.predict(user_input="Can you tell me about yourself.")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today? 

User: Can you tell me about yourself.[/INST][0m

[1m> Finished chain.[0m


"  Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today?"

In [None]:
llm_chain.predict(user_input="Today is Friday. What number day of the week is that?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today? 

User: Today is Friday. What number day of the week is that?[/INST][0m





[1m> Finished chain.[0m


'  AI: Great question! Friday is the 5th day of the week.'

In [None]:
llm_chain.predict(user_input="what is the day today?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week. 

User: what is the day today?[/INST][0m

[1m> Finished chain.[0m


'  AI: Today is Friday! üòä'

In [None]:
llm_chain.predict(user_input="What is my name?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI: Today is Friday! üòä 

User: What is my name?[/INST][0m

[1m> Finished chain.[0m


'  AI: Hello! My pleasure to assist you. Your name is Sam. üòä'

In [None]:
llm_chain.predict(user_input="Can you tell me about the olympics")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI: Today is Friday! üòä
Human: What is my name?
AI:   AI: Hello! My pleasure to assist you. Your name is Sam. üòä 

User: Can you tell me about the olympics[/INST][0m

[1m> Finished chain.[0m


"  Of course, I'd be happy to help! The Olympics are a major international sporting event that takes place every 4 years, featuring athletes from around the world competing in a variety of sports. The Olympic Games are overseen by the International Olympic Committee (IOC) and are held in a different city each time.\nThe Olympics include a wide range of sports, such as track and field, swimming, gymnastics, basketball, and many others. The events are held in various venues within the host city, and athletes compete for medals in each event. The country with the most medals at the end of the games is awarded the Olympic flag, and the host country is awarded the Olympic torch.\nThe Olympics have a rich history, dating back to 1896 in Athens, Greece. The games have evolved over the years, with new sports being added and others being removed. The Olympics have also been host to some of the most iconic moments in sports history, including Jesse Owens' four gold medals in Berlin in 1936 and M

In [None]:
llm_chain.predict(user_input="What have we talked about in this Chat?")



[1m> Entering new LLMChain chain...[0m
Prompt after formatting:
[32;1m[1;3m[INST]<<SYS>>
You are a helpful assistant, you always only answer for the assistant then you stop. read the chat history to get context
<</SYS>>

Chat History:

Human: Hi, my name is Sam
AI:   Hello Sam! It's nice to meet you. How can I assist you today?
Human: Can you tell me about yourself.
AI:   Of course! I'm just an AI designed to assist and provide helpful responses. I'm here to help you with any questions or tasks you may have. How can I assist you today?
Human: Today is Friday. What number day of the week is that?
AI:   AI: Great question! Friday is the 5th day of the week.
Human: what is the day today?
AI:   AI: Today is Friday! üòä
Human: What is my name?
AI:   AI: Hello! My pleasure to assist you. Your name is Sam. üòä
Human: Can you tell me about the olympics
AI:   Of course, I'd be happy to help! The Olympics are a major international sporting event that takes place every 4 years, featuring 




[1m> Finished chain.[0m


"  Of course! Here's what we've discussed in this chat:\n1. The assistant introduced themselves and asked how they could assist the user.\n2. The user asked the assistant to tell them about themselves, to which the assistant replied that they are just an AI designed to assist and provide helpful responses.\n3. The user asked what day of the week it is, and the assistant replied that it is Friday, which is the 5th day of the week.\n4. The user asked their name, and the assistant replied that their name is Sam.\n5. The user asked the assistant to tell them about the Olympics, and the assistant provided a detailed response about the Olympics, including their history, sports included, and the awarding of medals.\n6. Finally, the user asked the assistant what they have talked about in this chat, and the assistant replied that they have talked about the following topics:\n* The assistant's introduction and request to assist the user\n* The user's request to know more about themselves\n* The 