In [3]:
#!pip install llama-index chroma chromadb==0.3.25 pydantic==1.10.9 openai==0.27.8 bs4 tiktoken==0.4.0 langchain==0.0.235 huggingface_hub==0.16.4 sentence_transformers==2.2.2 pandas

In [54]:
import os
import requests
import bs4

from langchain.chat_models import ChatOpenAI
from langchain.vectorstores import Chroma
from langchain.text_splitter import RecursiveCharacterTextSplitter
from langchain.llms import OpenAI
from langchain.document_loaders import TextLoader
from langchain.chains import RetrievalQA
from langchain.embeddings import SentenceTransformerEmbeddings

In [55]:
model_name = "gpt-3.5-turbo"
llm = ChatOpenAI(model_name=model_name)

os.environ['OPENAI_API_KEY'] = "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"

In [21]:
def save_text_file(url, file_name):
    # Get the content of the file from the URL
    response = requests.get(url)
    
    if response.status_code == 200:
        with open(destination_path, 'w', encoding='utf-8') as file:
            file.write(response.text)
        return destination_path
    
    else:
        print(f"Failed to retrieve data. Status code: {response.status_code}")
        return None

In [22]:
url = "https://www.gutenberg.org/cache/epub/1998/pg1998.txt"
file_name = "thus_spake_zarathustra.txt"
current_directory = os.getcwd()
destination_path = os.path.join(current_directory, file_name)

saved_path = save_text_file(url, file_name)

In [45]:
loader = TextLoader(file_path=saved_path)
documents = loader.load()

In [46]:
len(documents)

1

In [47]:
from langchain.text_splitter import RecursiveCharacterTextSplitter

def split_docs(documents,chunk_size=1000,chunk_overlap=20):
  text_splitter = RecursiveCharacterTextSplitter(chunk_size=chunk_size, chunk_overlap=chunk_overlap)
  docs = text_splitter.split_documents(documents)
  return docs

docs = split_docs(documents)
print(len(docs))

776


In [48]:
from langchain.embeddings import SentenceTransformerEmbeddings
embeddings = SentenceTransformerEmbeddings(model_name="all-MiniLM-L6-v2")

In [49]:
from langchain.vectorstores import Chroma

persist_directory = "chroma_db"

vectordb = Chroma.from_documents(
    documents=docs, embedding=embeddings, persist_directory=persist_directory
)

vectordb.persist()

Batches:   0%|          | 0/25 [00:00<?, ?it/s]

In [51]:
from langchain.prompts import PromptTemplate                #prompt
from langchain.memory import ConversationBufferMemory       #memory
from langchain.chains import RetrievalQA                    #chain

In [60]:
template = """
You a professor who has indepth knowledge of 'Thus Spake Zarathustra'. Provide philosophical answers to the questions. 
Use quotes from the text to support your answers.

Context: {context}
-----------------------
History: {chat_history}
=======================
Human: {question}
Chatbot:

"""

# Create a prompt using the template
prompt = PromptTemplate(
    input_variables=["chat_history", "context", "question"],
    template=template
)

# Set up conversation memory
memory = ConversationBufferMemory(memory_key="chat_history", return_messages=True, input_key="question")

In [62]:
qa = RetrievalQA.from_chain_type(
    llm=llm,
    chain_type='stuff',
    retriever=vectordb.as_retriever(),
    verbose=True,
    chain_type_kwargs={
        "prompt": prompt,
        "memory": memory,
    }
)

In [57]:
# Function to print the response for a given query
def print_response_for_query(query):
    return print(qa.run({"query": query}))

In [61]:
query = "What are Nietzsche's views on morality?"
qa.run(query)



[1m> Entering new RetrievalQA chain...[0m


Batches:   0%|          | 0/1 [00:00<?, ?it/s]


[1m> Finished chain.[0m


"Nietzsche's views on morality are rooted in the idea that concepts of good and evil are merely means to an end, specifically, acquiring power. He argues that there are no absolute values of good and evil, but rather these are adopted by individuals and societies as a way to maintain their place in the world or to become dominant. Nietzsche distinguishes between two types of morality: the first, which he calls noble-morality, is creative and active, leading to an ascent in the line of life. The second, slave-morality, is passive and defensive, focused on self-preservation. Nietzsche advocates for noble-morality, seeing it as conducive to growth and vitality, while criticizing Christian moral values as belonging to a lower type of morality. He views morality as a tool for different types of individuals and societies to assert power and dominance."

In [59]:
query = "What are the most important ideas that 'Thus Spake Zarathustra' put forward?"
print_response_for_query(query)



[1m> Entering new RetrievalQA chain...[0m


Batches:   0%|          | 0/1 [00:00<?, ?it/s]


[1m> Finished chain.[0m
'Thus Spake Zarathustra' puts forward several important ideas. One of the central themes is the concept of the "Ubermensch" or the "Overman." Nietzsche presents the idea that humanity should strive to overcome its limitations and embrace a higher form of existence, transcending traditional moral and societal values. The Ubermensch is an individual who creates their own values and lives authentically, free from the constraints of herd mentality.

Another important idea in 'Thus Spake Zarathustra' is the notion of eternal recurrence. Nietzsche suggests that life is a cyclical process, and every moment will recur infinitely. This idea challenges individuals to embrace the present moment fully and live their lives in a way that they would be willing to repeat for eternity.

The book also explores the idea of the death of God and the consequences of living in a post-religious world. Nietzsche suggests that the absence of a divine authority creates a void that huma

In [63]:
query = "What life lessons does the book teaches?"
print_response_for_query(query)



[1m> Entering new RetrievalQA chain...[0m


Batches:   0%|          | 0/1 [00:00<?, ?it/s]


[1m> Finished chain.[0m
The book 'Thus Spake Zarathustra' teaches several life lessons. One of the key lessons is the importance of self-overcoming and personal growth. Zarathustra encourages individuals to strive for self-improvement and to constantly challenge themselves in order to reach higher levels of existence. As he says, "Man is a rope stretched between the animal and the Superman—a rope over an abyss. A dangerous crossing, a dangerous wayfaring, a dangerous looking-back, a dangerous trembling and halting" (Chapter LXXVI).

Another important lesson is the rejection of societal norms and the pursuit of individual freedom. Zarathustra urges individuals to break free from the herd mentality and to embrace their individuality. He states, "I beseech you, my brothers, remain true to the earth, and do not believe those who speak to you of otherworldly hopes! Poison-mixers are they, whether they know it or not" (Chapter LXXVII).

Furthermore, the book emphasizes the power of overco

In [10]:
# To cleanup, you can delete the collection
vectordb.delete_collection()
vectordb.persist()

# Or just nuke the persist directory
!rm -rf db/

Persisting DB to disk, putting it in the save folder db
