conversational-milo

Runtime error

App Files Files Community

vericudebuget commited on Apr 1, 2024

Commit

1e3869c

verified ·

1 Parent(s): 131d00b

Update app.py

Browse files

Files changed (1) hide show

app.py +16 -16

app.py CHANGED Viewed

@@ -1,9 +1,13 @@
-from huggingface_hub import InferenceClient
 import gradio as gr
 # Initialize the InferenceClient
 client = InferenceClient("mistralai/Mixtral-8x7B-Instruct-v0.1")
 def format_prompt(message, history):
     prompt = "<s>"
     for user_prompt, bot_response in history:
@@ -12,7 +16,7 @@ def format_prompt(message, history):
     prompt += f"[INST] {message} [/INST]"
     return prompt
-def generate(prompt, history, system_prompt, theme, temperature=0.9, max_new_tokens=9048, top_p=0.95, repetition_penalty=1.0):
     temperature = max(float(temperature), 1e-2)
     top_p = float(top_p)
@@ -34,23 +38,19 @@ def generate(prompt, history, system_prompt, theme, temperature=0.9, max_new_tok
         yield output
     return output
-def get_theme(theme_name):
-    theme_mapping = {
-        "Base": gr.themes.Base(),
-        "Default": gr.themes.Default(),
-        "Glass": gr.themes.Glass(),
-        "Monochrome": gr.themes.Monochrome(),
-        "Soft": gr.themes.Soft()
-    }
-    return theme_mapping.get(theme_name, gr.themes.Soft())
 additional_inputs = [
     gr.Textbox(label="System Prompt", max_lines=1, interactive=True),
     gr.Slider(label="Temperature", value=0.9, minimum=0.0, maximum=1.0, step=0.05, interactive=True, info="Higher values produce more diverse outputs"),
     gr.Slider(label="Max new tokens", value=9048, minimum=256, maximum=9048, step=64, interactive=True, info="The maximum numbers of new tokens"),
     gr.Slider(label="Top-p (nucleus sampling)", value=0.90, minimum=0.0, maximum=1, step=0.05, interactive=True, info="Higher values sample more low-probability tokens"),
-    gr.Slider(label="Repetition penalty", value=1.2, minimum=1.0, maximum=2.0, step=0.05, interactive=True, info="Penalize repeated tokens"),
-    gr.Dropdown(label="Theme", choices=["Default", "Base", "Glass", "Monochrome", "Soft"], interactive=True, info="Select a theme for the app")
 ]
 gr.ChatInterface(
@@ -58,7 +58,7 @@ gr.ChatInterface(
     chatbot=gr.Chatbot(show_label=True, show_share_button=True, show_copy_button=True, likeable=True, layout="panel"),
     additional_inputs=additional_inputs,
     title="ConvoLite",
-    description="Remember! The AI might give incorrect information about people, locations, history, etc...",
     concurrency_limit=20,
-    theme=get_theme(additional_inputs[-1].value)  # Use the selected theme
-).launch(show_api=False)

 import gradio as gr
+# Set the theme
+gr.Interface.load_theme(gr.themes.Soft())
 # Initialize the InferenceClient
 client = InferenceClient("mistralai/Mixtral-8x7B-Instruct-v0.1")
+# This code was mostly written by vericudebuget and gpt-4
 def format_prompt(message, history):
     prompt = "<s>"
     for user_prompt, bot_response in history:
     prompt += f"[INST] {message} [/INST]"
     return prompt
+def generate(prompt, history, system_prompt, temperature=0.9, max_new_tokens=9048, top_p=0.95, repetition_penalty=1.0):
     temperature = max(float(temperature), 1e-2)
     top_p = float(top_p)
         yield output
     return output
+    for response in stream:
+        output += response.token.text
+        if "http" in output:  # assuming the AI writes a direct image link in its response
+            yield {"image": output}  # Gradio will display the image
+        else:
+            yield output
 additional_inputs = [
     gr.Textbox(label="System Prompt", max_lines=1, interactive=True),
     gr.Slider(label="Temperature", value=0.9, minimum=0.0, maximum=1.0, step=0.05, interactive=True, info="Higher values produce more diverse outputs"),
     gr.Slider(label="Max new tokens", value=9048, minimum=256, maximum=9048, step=64, interactive=True, info="The maximum numbers of new tokens"),
     gr.Slider(label="Top-p (nucleus sampling)", value=0.90, minimum=0.0, maximum=1, step=0.05, interactive=True, info="Higher values sample more low-probability tokens"),
+    gr.Slider(label="Repetition penalty", value=1.2, minimum=1.0, maximum=2.0, step=0.05, interactive=True, info="Penalize repeated tokens")
 ]
 gr.ChatInterface(
     chatbot=gr.Chatbot(show_label=True, show_share_button=True, show_copy_button=True, likeable=True, layout="panel"),
     additional_inputs=additional_inputs,
     title="ConvoLite",
+    description= Remember! The AI might give incorect information about people, locations, history, etc...
     concurrency_limit=20,
+).launch(show_api=False,)