mixtral-46.7b-chat

Runtime error

App Files Files Community

vericudebuget commited on Apr 1

Commit

076fc13

•

1 Parent(s): 04b933e

Update app.py

Browse files

Files changed (1) hide show

app.py +17 -21

app.py CHANGED Viewed

@@ -12,7 +12,7 @@ def format_prompt(message, history):
     prompt += f"[INST] {message} [/INST]"
     return prompt
-def generate(prompt, history, system_prompt, theme, temperature=0.9, max_new_tokens=9048, top_p=0.95, repetition_penalty=1.0):
     temperature = max(float(temperature), 1e-2)
     top_p = float(top_p)
@@ -34,31 +34,27 @@ def generate(prompt, history, system_prompt, theme, temperature=0.9, max_new_tok
         yield output
     return output
-def get_theme(theme_name):
-    theme_mapping = {
-        "Base": gr.themes.Base(),
-        "Default": gr.themes.Default(),
-        "Glass": gr.themes.Glass(),
-        "Monochrome": gr.themes.Monochrome(),
-        "Soft": gr.themes.Soft()
-    }
-    return theme_mapping.get(theme_name, gr.themes.Default())
 additional_inputs = [
     gr.Textbox(label="System Prompt", max_lines=1, interactive=True),
     gr.Slider(label="Temperature", value=0.9, minimum=0.0, maximum=1.0, step=0.05, interactive=True, info="Higher values produce more diverse outputs"),
     gr.Slider(label="Max new tokens", value=9048, minimum=256, maximum=9048, step=64, interactive=True, info="The maximum numbers of new tokens"),
     gr.Slider(label="Top-p (nucleus sampling)", value=0.90, minimum=0.0, maximum=1, step=0.05, interactive=True, info="Higher values sample more low-probability tokens"),
-    gr.Slider(label="Repetition penalty", value=1.2, minimum=1.0, maximum=2.0, step=0.05, interactive=True, info="Penalize repeated tokens"),
-    gr.Dropdown(label="Theme", choices=["Default", "Base", "Glass", "Monochrome", "Soft"], interactive=True, info="Select a theme for the app")
 ]
-gr.ChatInterface(
-    fn=generate,
-    chatbot=gr.Chatbot(show_label=True, show_share_button=True, show_copy_button=True, likeable=True, layout="panel"),
-    additional_inputs=additional_inputs,
-    title="ConvoLite",
-    description="Remember! The AI might give incorrect information about people, locations, history, etc...",
-    concurrency_limit=20,
-    theme=get_theme(additional_inputs[-1].value)  # Use the selected theme
-).launch(show_api=False)

     prompt += f"[INST] {message} [/INST]"
     return prompt
+def generate(prompt, history, system_prompt, temperature=0.9, max_new_tokens=9048, top_p=0.95, repetition_penalty=1.0):
     temperature = max(float(temperature), 1e-2)
     top_p = float(top_p)
         yield output
     return output
+    for response in stream:
+        output += response.token.text
+        if "http" in output:  # assuming the AI writes a direct image link in its response
+            yield {"image": output}  # Gradio will display the image
+        else:
+            yield output
 additional_inputs = [
     gr.Textbox(label="System Prompt", max_lines=1, interactive=True),
     gr.Slider(label="Temperature", value=0.9, minimum=0.0, maximum=1.0, step=0.05, interactive=True, info="Higher values produce more diverse outputs"),
     gr.Slider(label="Max new tokens", value=9048, minimum=256, maximum=9048, step=64, interactive=True, info="The maximum numbers of new tokens"),
     gr.Slider(label="Top-p (nucleus sampling)", value=0.90, minimum=0.0, maximum=1, step=0.05, interactive=True, info="Higher values sample more low-probability tokens"),
+    gr.Slider(label="Repetition penalty", value=1.2, minimum=1.0, maximum=2.0, step=0.05, interactive=True, info="Penalize repeated tokens")
 ]
+with gr.themes.Soft():  # Apply the Soft theme
+    gr.ChatInterface(
+        fn=generate,
+        chatbot=gr.Chatbot(show_label=True, show_share_button=True, show_copy_button=True, likeable=True, layout="panel"),
+        additional_inputs=additional_inputs,
+        title="ConvoLite",
+        description= Remember! The AI might give incorrect information about people, locations, history, etc...
+        concurrency_limit=20,
+    ).launch(show_api=False,)