## Import Dependencies

In [1]:
from langchain.agents import ZeroShotAgent, Tool, AgentExecutor
from langchain import LLMChain
from langchain.tools import DuckDuckGoSearchRun
from langchain.llms import LlamaCpp
from langchain.callbacks.manager import CallbackManager
from langchain.callbacks.streaming_stdout import StreamingStdOutCallbackHandler

#Playwright
from langchain.agents.agent_toolkits import PlayWrightBrowserToolkit
from langchain.tools.playwright.utils import (
    create_async_playwright_browser,
    create_sync_playwright_browser,  # A synchronous browser is available, though it isn't compatible with jupyter.
)

#Python Tools
from langchain.agents.agent_toolkits import create_python_agent
from langchain.tools.python.tool import PythonREPLTool
from langchain.python import PythonREPL

# This import is required only for jupyter notebooks, since they have their own eventloop
import nest_asyncio

nest_asyncio.apply()

#Files ToolKit
from langchain.tools.file_management import (
    ReadFileTool,
    CopyFileTool,
    DeleteFileTool,
    MoveFileTool,
    WriteFileTool,
    ListDirectoryTool,
)
from langchain.agents.agent_toolkits import FileManagementToolkit
from tempfile import TemporaryDirectory

# We'll make a temporary directory to avoid clutter
working_directory = TemporaryDirectory()

## Load LLM

In [2]:
callback_manager = CallbackManager([StreamingStdOutCallbackHandler()])
model_path = "../../../../../Models/GGML/llama-2-7b-chat.ggmlv3.q8_0.bin"
llm = LlamaCpp(
    model_path=model_path,
    n_ctx=2048,
    f16_kv=True,  # MUST set to True, otherwise you will run into problem after a couple of calls
    callback_manager=callback_manager,
    verbose=True,
)

llama.cpp: loading model from ../../../../../Models/GGML/llama-2-7b-chat.ggmlv3.q8_0.bin
llama_model_load_internal: format     = ggjt v3 (latest)
llama_model_load_internal: n_vocab    = 32000
llama_model_load_internal: n_ctx      = 2048
llama_model_load_internal: n_embd     = 4096
llama_model_load_internal: n_mult     = 256
llama_model_load_internal: n_head     = 32
llama_model_load_internal: n_head_kv  = 32
llama_model_load_internal: n_layer    = 32
llama_model_load_internal: n_rot      = 128
llama_model_load_internal: n_gqa      = 1
llama_model_load_internal: rnorm_eps  = 5.0e-06
llama_model_load_internal: n_ff       = 11008
llama_model_load_internal: freq_base  = 10000.0
llama_model_load_internal: freq_scale = 1
llama_model_load_internal: ftype      = 7 (mostly Q8_0)
llama_model_load_internal: model size = 7B
llama_model_load_internal: ggml ctx size =    0.08 MB
llama_model_load_internal: mem required  = 6828.73 MB (+ 1024.00 MB per state)
llama_new_context_with_model: kv self size 

## Load Tools

In [3]:
async_browser = create_async_playwright_browser()
scrape_toolkit = PlayWrightBrowserToolkit.from_browser(async_browser=async_browser)
scraping_tools = scrape_toolkit.get_tools()

In [4]:
tools_by_name = {tool.name: tool for tool in scraping_tools}
print(tools_by_name.keys())
navigate_tool = tools_by_name["navigate_browser"]
extract_text = tools_by_name["extract_text"]

dict_keys(['click_element', 'navigate_browser', 'previous_webpage', 'extract_text', 'extract_hyperlinks', 'get_elements', 'current_webpage'])


In [5]:
files_tools = FileManagementToolkit(
    root_dir=str(working_directory.name),
    selected_tools=["read_file", "write_file", "list_directory"],
).get_tools()
read_tool, write_tool, list_tool = files_tools

In [6]:
DuckDuckGoSearch = DuckDuckGoSearchRun()
PythonTool =PythonREPLTool()

tools = [
    Tool(
        name="DuckDuckGo",
        func=DuckDuckGoSearch.run,
        description="useful for when you need to answer questions about current events",
    ),
    Tool(
        name="Read directory",
        func=read_tool.run,
        description="Read files within a directory",
    ),
    Tool(
        name="Write directory",
        func=write_tool.run,
        description="Write files to a directory",
    ),
    Tool(
        name="Extract text from webpage",
        func=extract_text.run,
        description="Extracts text from a webpage",
    ),
    extract_text,
    navigate_tool, PythonTool, 
]

In [7]:
prefix = """Answer the following questions as best you You have access to the following tools:"""
suffix = """Begin! Remember to be as creative as you can"

Question: {input}
{agent_scratchpad}"""

prompt = ZeroShotAgent.create_prompt(
    tools, prefix=prefix, suffix=suffix, input_variables=["input", "agent_scratchpad"]
)

In [8]:
llm_chain = LLMChain(llm=llm, prompt=prompt)

In [9]:
tool_names = [tool.name for tool in tools]
agent = ZeroShotAgent(llm_chain=llm_chain, allowed_tools=tool_names)

In [10]:
agent_executor = AgentExecutor.from_agent_and_tools(
    agent=agent, tools=tools, verbose=True, handle_parsing_errors=True
)

In [11]:
agent_executor.run("""
Let's play a game. You're software engineer with godlike programming skills. Help me make a voice assistant that will basically be the voice of my laptop. I want to call him "Bryan". We also would like to use a voice installed locally on my machine (windows). The voice assistant must be in python. When i press the spacebar it must start recording and stop when I press the spacebar again. Then Bryan must speak what I just said. Use the Langchain code below to engineer the voice assistant code.

LiteLLM Code:
from litellm import completion

os.environ["COHERE_API_KEY"] = "cohere key"

messages = [{ "content": "Hello, how are you?","role": "user"}]
# cohere call
response = completion(model="command-nightly", messages)
print(response.choices[0].message.content)

In the above code the response is what I'd like for Bryan to speak. The content in messages is what I'd like to query Bryan with.""")



[1m> Entering new AgentExecutor chain...[0m
Thought: What should we do first?


llama_print_timings:        load time = 189757.29 ms
llama_print_timings:      sample time =   143.79 ms /    10 runs   (   14.38 ms per token,    69.55 tokens per second)
llama_print_timings: prompt eval time = 864828.21 ms /   534 tokens ( 1619.53 ms per token,     0.62 tokens per second)
llama_print_timings:        eval time = 1136898.11 ms /     9 runs   (126322.01 ms per token,     0.01 tokens per second)
llama_print_timings:       total time = 2010180.99 ms


[32;1m[1;3mThought: What should we do first?[0m
Observation: Invalid Format: Missing 'Action:' after 'Thought:
Thought:

Llama.generate: prefix-match hit


 We can use the Read directory action to read the file containing the voices installed locally on my machine (windows).
Action Input: Read directory "C:\Program Files\Bryan Voice"
Action