Code
cookbook/02_examples/teams/multimodal/image_to_text.py
Usage
1
Set up your virtual environment
2
Install required libraries
3
Set environment variables
4
Add sample image
5
Run the agent
Documentation Index
Fetch the complete documentation index at: /llms.txt
Use this file to discover all available pages before exploring further.
from pathlib import Path
from agno.agent import Agent
from agno.media import Image
from agno.models.openai import OpenAIResponses
from agno.team import Team
image_analyzer = Agent(
name="Image Analyst",
role="Analyze and describe images in detail",
model=OpenAIResponses(id="gpt-5.2"),
instructions=[
"Analyze images carefully and provide detailed descriptions",
"Focus on visual elements, composition, and key details",
],
)
creative_writer = Agent(
name="Creative Writer",
role="Create engaging stories and narratives",
model=OpenAIResponses(id="gpt-5.2"),
instructions=[
"Transform image descriptions into compelling fiction stories",
"Use vivid language and creative storytelling techniques",
],
)
# Create a team for collaborative image-to-text processing
image_team = Team(
name="Image Story Team",
model=OpenAIResponses(id="gpt-5.2"),
members=[image_analyzer, creative_writer],
instructions=[
"Work together to create compelling fiction stories from images.",
"Image Analyst: First analyze the image for visual details and context.",
"Creative Writer: Transform the analysis into engaging fiction narratives.",
"Ensure the story captures the essence and mood of the image.",
],
markdown=True,
)
image_path = Path(__file__).parent.joinpath("sample.jpg")
image_team.print_response(
"Write a 3 sentence fiction story about the image",
images=[Image(filepath=image_path)],
)
Set up your virtual environment
uv venv --python 3.12
source .venv/bin/activate
uv venv --python 3.12
.venv\Scripts\activate
Install required libraries
uv pip install agno
Set environment variables
export OPENAI_API_KEY=****
Add sample image
# Add a sample.jpg image file in the same directory as the script
Run the agent
python cookbook/02_examples/teams/multimodal/image_to_text.py
Was this page helpful?