In [2]:
"""
Installs the `arxiv` Python library, which provides a simple interface for accessing the arXiv.org API.

The `arxiv` library allows you to search for and retrieve metadata about research papers published on arXiv.org, a popular open-access repository for scientific papers.
"""
!pip install arxiv

Collecting arxiv
  Downloading arxiv-2.1.3-py3-none-any.whl.metadata (6.1 kB)
Collecting feedparser~=6.0.10 (from arxiv)
  Downloading feedparser-6.0.11-py3-none-any.whl.metadata (2.4 kB)
Collecting sgmllib3k (from feedparser~=6.0.10->arxiv)
  Downloading sgmllib3k-1.0.0.tar.gz (5.8 kB)
  Preparing metadata (setup.py) ... [?25l[?25hdone
Downloading arxiv-2.1.3-py3-none-any.whl (11 kB)
Downloading feedparser-6.0.11-py3-none-any.whl (81 kB)
[2K   [90m━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━━[0m [32m81.3/81.3 kB[0m [31m6.2 MB/s[0m eta [36m0:00:00[0m
[?25hBuilding wheels for collected packages: sgmllib3k
  Building wheel for sgmllib3k (setup.py) ... [?25l[?25hdone
  Created wheel for sgmllib3k: filename=sgmllib3k-1.0.0-py3-none-any.whl size=6047 sha256=ef9df83cc2452e8cf284961f52cb2b917efa7dcf93527d605c8c821609efdf07
  Stored in directory: /root/.cache/pip/wheels/f0/69/93/a47e9d621be168e9e33c7ce60524393c0b92ae83cf6c6e89c5
Successfully built sgmllib3k
Installing collected packag

In [3]:
"""
faiss-cpu: Fast similarity search and clustering for large datasets.
langchain: Framework for building LLM-powered applications.
langchain_community: Community-driven extensions for LangChain.
langchain_openai: Using OpenAI LLMs with LangChain.
openai: Accessing OpenAI's API for various machine learning models.
"""

!pip install wikipedia faiss-cpu langchain langchain_community langchain_openai openai

Collecting wikipedia
  Downloading wikipedia-1.4.0.tar.gz (27 kB)
  Preparing metadata (setup.py) ... [?25l[?25hdone
Collecting faiss-cpu
  Downloading faiss_cpu-1.8.0.post1-cp310-cp310-manylinux_2_17_x86_64.manylinux2014_x86_64.whl.metadata (3.7 kB)
Collecting langchain
  Downloading langchain-0.3.1-py3-none-any.whl.metadata (7.1 kB)
Collecting langchain_community
  Downloading langchain_community-0.3.1-py3-none-any.whl.metadata (2.8 kB)
Collecting langchain_openai
  Downloading langchain_openai-0.2.0-py3-none-any.whl.metadata (2.6 kB)
Collecting openai
  Downloading openai-1.48.0-py3-none-any.whl.metadata (24 kB)
Collecting langchain-core<0.4.0,>=0.3.6 (from langchain)
  Downloading langchain_core-0.3.6-py3-none-any.whl.metadata (6.3 kB)
Collecting langchain-text-splitters<0.4.0,>=0.3.0 (from langchain)
  Downloading langchain_text_splitters-0.3.0-py3-none-any.whl.metadata (2.3 kB)
Collecting langsmith<0.2.0,>=0.1.17 (from langchain)
  Downloading langsmith-0.1.128-py3-none-any.whl

In [11]:
"""
Loads the OpenAI API key from the environment and configures the ChatOpenAI language model with the specified parameters.

The `load_dotenv()` function loads environment variables from a `.env` file, which is used to store sensitive API keys and other configuration values. The `os.environ['OPENAI_API_KEY']` line sets the `OPENAI_API_KEY` environment variable to the value retrieved from the `.env` file.

The `ChatOpenAI` class from the `langchain_openai` module is then used to create an instance of the OpenAI language model, with the specified model name and temperature parameter. This instance can be used to generate text using the OpenAI API.
"""
from dotenv import load_dotenv
import getpass
load_dotenv()
import os
os.environ['OPENAI_API_KEY']=getpass.getpass()
from langchain_openai import ChatOpenAI

llm = ChatOpenAI(model="gpt-3.5-turbo-0125", temperature=0)

··········


In [5]:
"""
This module provides two utility classes for interacting with Wikipedia:

`WikipediaQueryRun`: A class that allows you to run a Wikipedia query and retrieve the results.
`WikipediaAPIWrapper`: A wrapper around the Wikipedia API that provides a more user-friendly interface for querying and retrieving information from Wikipedia.

These utilities can be used to incorporate Wikipedia-based functionality into your applications.
"""
from langchain_community.tools import WikipediaQueryRun
from langchain_community.utilities import WikipediaAPIWrapper

In [6]:
"""
Initializes a WikipediaAPIWrapper instance with a maximum of 1 top result and a maximum of 200 characters per document content.
Then, creates a WikipediaQueryRun instance using the configured WikipediaAPIWrapper.
"""
api_wrapper=WikipediaAPIWrapper(top_k_results=1,doc_content_chars_max=200)
wiki=WikipediaQueryRun(api_wrapper=api_wrapper)

In [7]:
wiki.name

'wikipedia'

In [12]:
"""
This code sets up a document retrieval system using the LangChain library. It loads web-based documents from a specified URL, splits the documents into smaller chunks, and creates a FAISS vector store using OpenAI embeddings. The resulting retriever can be used to search and retrieve relevant documents based on user queries.

The key steps are:
1. Load web-based documents using the WebBaseLoader
2. Split the documents into smaller chunks using the RecursiveCharacterTextSplitter
3. Create a FAISS vector store from the document chunks using OpenAI embeddings
4. Obtain a retriever from the vector store that can be used to search and retrieve relevant documents
"""
from langchain_community.document_loaders import WebBaseLoader
from langchain_community.vectorstores import FAISS
from langchain_openai import OpenAIEmbeddings
from langchain_text_splitters import RecursiveCharacterTextSplitter

loader=WebBaseLoader("https://docs.smith.langchain.com/")
docs=loader.load()
documents=RecursiveCharacterTextSplitter(chunk_size=1000,chunk_overlap=200).split_documents(docs)
vectordb=FAISS.from_documents(documents,OpenAIEmbeddings())
retriever=vectordb.as_retriever()
retriever


VectorStoreRetriever(tags=['FAISS', 'OpenAIEmbeddings'], vectorstore=<langchain_community.vectorstores.faiss.FAISS object at 0x7a65a0c147f0>, search_kwargs={})

In [13]:
"""
Creates a LangChain tool that allows searching for information about LangSmith.

This tool should be used for any questions about LangSmith, as it provides a consistent and reliable way to search for relevant information.

Args:
    retriever (Retriever): The retriever to use for searching.
    name (str): The name of the tool.
    description (str): A description of the tool's purpose.

Returns:
    Tool: A LangChain tool that can be used to search for information about LangSmith.
"""
from langchain.tools.retriever import create_retriever_tool
retriever_tool=create_retriever_tool(retriever,"langsmith_search",
                      "Search for information about LangSmith. For any questions about LangSmith, you must use this tool!")

In [14]:
retriever_tool.name

'langsmith_search'

In [15]:
"""
Provides an ArxivAPIWrapper and ArxivQueryRun tool for interacting with the Arxiv API.

The ArxivAPIWrapper is a wrapper around the Arxiv API that allows for querying the API and returning a limited number of results with a maximum character count for the content.

The ArxivQueryRun tool uses the ArxivAPIWrapper to run a query against the Arxiv API and return the results.
"""
## Arxiv Tool
from langchain_community.utilities import ArxivAPIWrapper
from langchain_community.tools import ArxivQueryRun

arxiv_wrapper=ArxivAPIWrapper(top_k_results=1, doc_content_chars_max=200)
arxiv=ArxivQueryRun(api_wrapper=arxiv_wrapper)
arxiv.name

'arxiv'

In [26]:
#Mathematics calculator tool
from langchain.chains import LLMMathChain
from langchain.agents import Tool
problem_chain = LLMMathChain.from_llm(llm=llm)
math_tool = Tool.from_function(name="Calculator",
                func=problem_chain.run,
                 description="Useful for when you need to answer questions about math. This tool is only for math questions and nothing else. Only input math expressions.")

In [17]:
tools=[wiki,arxiv,retriever_tool,math_tool]

In [18]:
tools

[WikipediaQueryRun(api_wrapper=WikipediaAPIWrapper(wiki_client=<module 'wikipedia' from '/usr/local/lib/python3.10/dist-packages/wikipedia/__init__.py'>, top_k_results=1, lang='en', load_all_available_meta=False, doc_content_chars_max=200)),
 ArxivQueryRun(api_wrapper=ArxivAPIWrapper(arxiv_search=<class 'arxiv.Search'>, arxiv_exceptions=(<class 'arxiv.ArxivError'>, <class 'arxiv.UnexpectedEmptyPageError'>, <class 'arxiv.HTTPError'>), top_k_results=1, ARXIV_MAX_QUERY_LENGTH=300, continue_on_failure=False, load_max_docs=100, load_all_available_meta=False, doc_content_chars_max=200)),
 Tool(name='langsmith_search', description='Search for information about LangSmith. For any questions about LangSmith, you must use this tool!', args_schema=<class 'langchain_core.tools.retriever.RetrieverInput'>, func=functools.partial(<function _get_relevant_documents at 0x7a65a2427f40>, retriever=VectorStoreRetriever(tags=['FAISS', 'OpenAIEmbeddings'], vectorstore=<langchain_community.vectorstores.faiss.F

In [19]:
"""
Retrieves a prompt from the LangChain hub, which is a collection of pre-built prompts and agents for language models.

The `hub.pull()` function is used to fetch the prompt with the identifier "hwchase17/openai-functions-agent". This prompt can be used as a starting point for building an agent that interacts with OpenAI's language model functions.

The retrieved prompt is stored in the `prompt` variable, which can then be used to configure and customize the agent's behavior.
"""
from langchain import hub
# Get the prompt to use - you can modify this!
prompt = hub.pull("hwchase17/openai-functions-agent")
prompt.messages



[SystemMessagePromptTemplate(prompt=PromptTemplate(input_variables=[], input_types={}, partial_variables={}, template='You are a helpful assistant'), additional_kwargs={}),
 MessagesPlaceholder(variable_name='chat_history', optional=True),
 HumanMessagePromptTemplate(prompt=PromptTemplate(input_variables=['input'], input_types={}, partial_variables={}, template='{input}'), additional_kwargs={}),
 MessagesPlaceholder(variable_name='agent_scratchpad')]

In [20]:
"""
Creates an OpenAI tools agent using the provided language model (llm), tools, and prompt.

The agent is a powerful tool that can leverage various external tools and services to accomplish complex tasks. It uses the provided language model to understand natural language instructions, and then selects and invokes the appropriate tools to carry out those instructions.

Args:
    llm (LLMChain): The language model to use for the agent.
    tools (list): A list of tools that the agent can use to accomplish tasks.
    prompt (PromptTemplate): The prompt template to use for the agent.

Returns:
    Agent: The created OpenAI tools agent.
"""
### Agents
from langchain.agents import create_openai_tools_agent
agent=create_openai_tools_agent(llm,tools,prompt)

In [21]:
## Agent Executer
from langchain.agents import AgentExecutor
agent_executor=AgentExecutor(agent=agent,tools=tools,verbose=True)
agent_executor

AgentExecutor(verbose=True, agent=RunnableMultiActionAgent(runnable=RunnableAssign(mapper={
  agent_scratchpad: RunnableLambda(lambda x: format_to_openai_tool_messages(x['intermediate_steps']))
})
| ChatPromptTemplate(input_variables=['agent_scratchpad', 'input'], optional_variables=['chat_history'], input_types={'chat_history': list[typing.Annotated[typing.Union[typing.Annotated[langchain_core.messages.ai.AIMessage, Tag(tag='ai')], typing.Annotated[langchain_core.messages.human.HumanMessage, Tag(tag='human')], typing.Annotated[langchain_core.messages.chat.ChatMessage, Tag(tag='chat')], typing.Annotated[langchain_core.messages.system.SystemMessage, Tag(tag='system')], typing.Annotated[langchain_core.messages.function.FunctionMessage, Tag(tag='function')], typing.Annotated[langchain_core.messages.tool.ToolMessage, Tag(tag='tool')], typing.Annotated[langchain_core.messages.ai.AIMessageChunk, Tag(tag='AIMessageChunk')], typing.Annotated[langchain_core.messages.human.HumanMessageChunk, Tag

In [22]:
agent_executor.invoke({"input":"Tell me about Langsmith"})



[1m> Entering new AgentExecutor chain...[0m
[32;1m[1;3m
Invoking: `langsmith_search` with `{'query': 'Langsmith'}`


[0m[38;5;200m[1;3mGet started with LangSmith | ü¶úÔ∏èüõ†Ô∏è LangSmith

Skip to main contentGo to API DocsSearchRegionUSEUGo to AppQuick startTutorialsHow-to guidesConceptsReferencePricingSelf-hostingLangGraph CloudQuick startOn this pageGet started with LangSmithLangSmith is a platform for building production-grade LLM applications. It allows you to closely monitor and evaluate your application, so you can ship quickly and with confidence. Use of LangChain is not necessary - LangSmith works on its own!1. Install LangSmith‚ÄãPythonTypeScriptpip install -U langsmithyarn add langsmith2. Create an API key‚ÄãTo create an API key head to the Settings page. Then click Create API Key.3. Set up your environment‚ÄãShellexport LANGCHAIN_TRACING_V2=trueexport LANGCHAIN_API_KEY=<your-api-key># The below examples use the OpenAI API, though it's not necessary in generalexpor

{'input': 'Tell me about Langsmith',
 'output': 'LangSmith is a platform for building production-grade LLM (Large Language Model) applications. It allows you to closely monitor and evaluate your application, enabling you to ship quickly and with confidence. LangSmith works independently, and the use of LangChain is not necessary for its operation.\n\nTo get started with LangSmith, you can follow these steps:\n1. Install LangSmith using pip for Python or yarn for TypeScript:\n   - `pip install -U langsmith`\n   - `yarn add langsmith`\n\n2. Create an API key by heading to the Settings page and clicking on Create API Key.\n\n3. Set up your environment by exporting the necessary variables like `LANGCHAIN_TRACING_V2`, `LANGCHAIN_API_KEY`, and `OPENAI_API_KEY`.\n\n4. Log your first trace to LangSmith for LangChain users. You can use the LangSmith SDK directly or follow the tracing guide outlined for LangChain users.\n\nLangSmith provides multiple ways to log traces, and you can define datase

In [23]:
agent_executor.invoke({"input":"What's the paper 1605.08386 about?"})



[1m> Entering new AgentExecutor chain...[0m
[32;1m[1;3m
Invoking: `arxiv` with `{'query': '1605.08386'}`


[0m[33;1m[1;3mPublished: 2016-05-26
Title: Heat-bath random walks with Markov bases
Authors: Caprice Stanley, Tobias Windisch
Summary: Graphs on lattice points are studied whose edges come from a finite set of
allo[0m[32;1m[1;3mThe paper with the identifier 1605.08386 is titled "Heat-bath random walks with Markov bases" by Caprice Stanley and Tobias Windisch. The paper discusses the study of graphs on lattice points where the edges come from a finite set of allocations.[0m

[1m> Finished chain.[0m


{'input': "What's the paper 1605.08386 about?",
 'output': 'The paper with the identifier 1605.08386 is titled "Heat-bath random walks with Markov bases" by Caprice Stanley and Tobias Windisch. The paper discusses the study of graphs on lattice points where the edges come from a finite set of allocations.'}

In [24]:
agent_executor.invoke({"input":"tell me about virat kohli"})



[1m> Entering new AgentExecutor chain...[0m
[32;1m[1;3m
Invoking: `wikipedia` with `{'query': 'Virat Kohli'}`


[0m[36;1m[1;3mPage: Virat Kohli
Summary: Virat Kohli (Hindi pronunciation: [ʋɪˈɾɑːʈ ˈkoːɦli] ; born 5 November 1988) is an Indian international cricketer who plays Test and ODI cricket for the Indian national team.[0m[32;1m[1;3mVirat Kohli is an Indian international cricketer who plays Test and ODI cricket for the Indian national team. He was born on 5 November 1988. If you would like more information about him, feel free to ask![0m

[1m> Finished chain.[0m


{'input': 'tell me about virat kohli',
 'output': 'Virat Kohli is an Indian international cricketer who plays Test and ODI cricket for the Indian national team. He was born on 5 November 1988. If you would like more information about him, feel free to ask!'}

In [25]:
agent_executor.invoke({"input":"what is 35*4+5/2"})



[1m> Entering new AgentExecutor chain...[0m
[32;1m[1;3m
Invoking: `Calculator` with `35*4+5/2`


[0m[36;1m[1;3mAnswer: 142.5[0m[32;1m[1;3mThe result of the expression 35*4+5/2 is 142.5.[0m

[1m> Finished chain.[0m


{'input': 'what is 35*4+5/2',
 'output': 'The result of the expression 35*4+5/2 is 142.5.'}