Spaces:

JeCabrera
/

copywriter

Running

App Files Files Community

JeCabrera commited on Jan 17

Commit

69a8ba9

verified ·

1 Parent(s): 88c5efb

Update app.py

Browse files

Files changed (1) hide show

app.py +150 -86

app.py CHANGED Viewed

@@ -1,98 +1,162 @@
-import gradio as gr
 import os
 import time
 import google.generativeai as genai
-from mimetypes import MimeTypes
-# Configurar la API de Gemini
-genai.configure(api_key=os.environ["GOOGLE_API_KEY"])
-def upload_and_process_file(file):
-    """
-    Sube y procesa un archivo para usarlo con el modelo de Gemini.
-    """
-    if isinstance(file, gr.inputs.File):
-        file_path = file.name  # Obtener la ruta temporal del archivo
-    else:
-        raise ValueError("El archivo no se subió correctamente.")
-    # Detectar el tipo MIME del archivo
-    mime = MimeTypes()
-    mime_type, _ = mime.guess_type(file_path)
-    if not mime_type:
-        raise ValueError("No se pudo determinar el tipo MIME del archivo.")
-    # Subir el archivo a Gemini
-    print(f"Subiendo el archivo '{file_path}' con MIME type '{mime_type}'...")
-    file = genai.upload_file(file_path, mime_type=mime_type)
-    print(f"Archivo subido: {file.display_name}, URI: {file.uri}")
-    # Esperar a que el archivo esté activo
-    wait_for_files_active([file])
-    return file
-def wait_for_files_active(files):
-    """
-    Espera a que los archivos subidos a Gemini estén activos y listos para su uso.
-    """
-    print("Esperando el procesamiento de los archivos...")
     for file in files:
-        status = genai.get_file(file.name)
-        while status.state.name == "PROCESSING":
-            print(".", end="", flush=True)
-            time.sleep(5)
-            status = genai.get_file(file.name)
-        if status.state.name != "ACTIVE":
-            raise Exception(f"El archivo {file.name} no pudo procesarse correctamente.")
-    print("\nTodos los archivos están listos.")
-# Configuración del modelo generativo
-generation_config = {
-    "temperature": 1,
-    "top_p": 0.95,
-    "top_k": 40,
-    "max_output_tokens": 8192,
-    "response_mime_type": "text/plain",
-}
-model = genai.GenerativeModel(
-    model_name="gemini-1.5-flash",
-    generation_config=generation_config,
-)
-def start_chat_with_file(file, user_input):
-    """
-    Inicia una conversación con el modelo utilizando un archivo como entrada.
-    """
-    chat_session = model.start_chat(
-        history=[
-            {
-                "role": "user",
-                "parts": [file],
-            },
-        ]
     )
-    # Enviar mensaje al modelo
-    response = chat_session.send_message(user_input)
-    return response.text
-# Interfaz de Gradio
-def process_file(file, user_input):
-    processed_file = upload_and_process_file(file)
-    response = start_chat_with_file(processed_file, user_input)
-    return response
-# Crear la interfaz de Gradio
-iface = gr.Interface(
-    fn=process_file,
-    inputs=[gr.File(label="Sube tu archivo"), gr.Textbox(label="Escribe tu mensaje")],
-    outputs="text",
-    live=True
 )
-# Ejecutar la interfaz
-iface.launch()

+TITLE = """<h1 align="center">Gemini Playground ✨</h1>"""
+SUBTITLE = """<h2 align="center">Play with Gemini Pro and Gemini Pro Vision</h2>"""
 import os
 import time
+import uuid
+from typing import List, Tuple, Optional, Union
 import google.generativeai as genai
+import gradio as gr
+from PIL import Image
+from dotenv import load_dotenv
+# Cargar las variables de entorno desde el archivo .env
+load_dotenv()
+print("google-generativeai:", genai.__version__)
+# Obtener la clave de la API de las variables de entorno
+GOOGLE_API_KEY = os.getenv("GOOGLE_API_KEY")
+# Verificar que la clave de la API esté configurada
+if not GOOGLE_API_KEY:
+    raise ValueError("GOOGLE_API_KEY is not set in environment variables.")
+IMAGE_CACHE_DIRECTORY = "/tmp"
+IMAGE_WIDTH = 512
+CHAT_HISTORY = List[Tuple[Optional[Union[Tuple[str], str]], Optional[str]]]
+def preprocess_image(image: Image.Image) -> Optional[Image.Image]:
+    if image:
+        image_height = int(image.height * IMAGE_WIDTH / image.width)
+        return image.resize((IMAGE_WIDTH, image_height))
+def cache_pil_image(image: Image.Image) -> str:
+    image_filename = f"{uuid.uuid4()}.jpeg"
+    os.makedirs(IMAGE_CACHE_DIRECTORY, exist_ok=True)
+    image_path = os.path.join(IMAGE_CACHE_DIRECTORY, image_filename)
+    image.save(image_path, "JPEG")
+    return image_path
+def upload(files: Optional[List[str]], chatbot: CHAT_HISTORY) -> CHAT_HISTORY:
     for file in files:
+        image = Image.open(file).convert('RGB')
+        image_preview = preprocess_image(image)
+        if image_preview:
+            gr.Image(image_preview).render()
+        image_path = cache_pil_image(image)
+        chatbot.append(((image_path,), None))
+    return chatbot
+def user(text_prompt: str, chatbot: CHAT_HISTORY):
+    if text_prompt:
+        chatbot.append((text_prompt, None))
+    return "", chatbot
+def bot(
+    files: Optional[List[str]],
+    model_choice: str,
+    system_instruction: Optional[str],  # Sistema de instrucciones opcional
+    chatbot: CHAT_HISTORY
+):
+    if not GOOGLE_API_KEY:
+        raise ValueError("GOOGLE_API_KEY is not set.")
+    genai.configure(api_key=GOOGLE_API_KEY)
+    generation_config = genai.types.GenerationConfig(
+        temperature=0.7,
+        max_output_tokens=8192,
+        top_k=10,
+        top_p=0.9
+    )
+    # Usar el valor por defecto para system_instruction si está vacío
+    if not system_instruction:
+        system_instruction = "1"  # O puedes poner un valor predeterminado como "No system instruction provided."
+    text_prompt = [chatbot[-1][0]] if chatbot and chatbot[-1][0] and isinstance(chatbot[-1][0], str) else []
+    image_prompt = [preprocess_image(Image.open(file).convert('RGB')) for file in files] if files else []
+    model = genai.GenerativeModel(
+        model_name=model_choice,
+        generation_config=generation_config,
+        system_instruction=system_instruction  # Usar el valor por defecto si está vacío
     )
+    response = model.generate_content(text_prompt + image_prompt, stream=True, generation_config=generation_config)
+    chatbot[-1][1] = ""
+    for chunk in response:
+        for i in range(0, len(chunk.text), 10):
+            section = chunk.text[i:i + 10]
+            chatbot[-1][1] += section
+            time.sleep(0.01)
+            yield chatbot
+# Componente para el acordeón que contiene el cuadro de texto para la instrucción del sistema
+system_instruction_component = gr.Textbox(
+    placeholder="Enter system instruction...",
+    show_label=True,
+    scale=8
+)
+# Definir los componentes de entrada y salida
+chatbot_component = gr.Chatbot(label='Gemini', bubble_full_width=False, scale=2, height=300)
+text_prompt_component = gr.Textbox(placeholder="Message...", show_label=False, autofocus=True, scale=8)
+upload_button_component = gr.UploadButton(label="Upload Images", file_count="multiple", file_types=["image"], scale=1)
+run_button_component = gr.Button(value="Run", variant="primary", scale=1)
+model_choice_component = gr.Dropdown(
+    choices=["gemini-1.5-flash", "gemini-2.0-flash-exp", "gemini-1.5-pro"],
+    value="gemini-1.5-flash",
+    label="Select Model",
+    scale=2
 )
+user_inputs = [text_prompt_component, chatbot_component]
+bot_inputs = [upload_button_component, model_choice_component, system_instruction_component, chatbot_component]
+# Definir la interfaz de usuario
+with gr.Blocks() as demo:
+    gr.HTML(TITLE)
+    gr.HTML(SUBTITLE)
+    with gr.Column():
+        # Campo de selección de modelo arriba
+        model_choice_component.render()
+        chatbot_component.render()
+        with gr.Row():
+            text_prompt_component.render()
+            upload_button_component.render()
+            run_button_component.render()
+        # Crear el acordeón para la instrucción del sistema al final
+        with gr.Accordion("System Instruction", open=False):  # Acordeón cerrado por defecto
+            system_instruction_component.render()
+    run_button_component.click(
+        fn=user,
+        inputs=user_inputs,
+        outputs=[text_prompt_component, chatbot_component],
+        queue=False
+    ).then(
+        fn=bot, inputs=bot_inputs, outputs=[chatbot_component],
+    )
+    text_prompt_component.submit(
+        fn=user,
+        inputs=user_inputs,
+        outputs=[text_prompt_component, chatbot_component],
+        queue=False
+    ).then(
+        fn=bot, inputs=bot_inputs, outputs=[chatbot_component],
+    )
+    upload_button_component.upload(
+        fn=upload,
+        inputs=[upload_button_component, chatbot_component],
+        outputs=[chatbot_component],
+        queue=False
+    )
+# Lanzar la aplicación
+demo.queue(max_size=99).launch(debug=False, show_error=True)