{"component":{"type":17,"accent_color":7742708,"components":[{"type":9,"components":[{"type":10,"content":"## [Llama 3.2 11B Vision Instruct](https://openrouter.ai/meta-llama/llama-3.2-11b-vision-instruct)\nmeta-llama · 131K context"},{"type":10,"content":"Llama 3.2 11B Vision is a multimodal model with 11 billion parameters, designed to handle tasks combining visual and textual data. It excels in tasks such as image captioning and visual question answering, bridging the gap between language generation and visual reasoning.…"}],"accessory":{"type":11,"media":{"url":"https://openrouter.ai/images/icons/Meta.png"}}},{"type":14,"divider":true,"spacing":1},{"type":1,"components":[{"type":2,"style":5,"label":"Chat","url":"https://openrouter.ai/chat?models=meta-llama%2Fllama-3.2-11b-vision-instruct","emoji":{"name":"💬"}},{"type":2,"style":5,"label":"Compare","url":"https://openrouter.ai/compare/meta-llama/llama-3.2-11b-vision-instruct"},{"type":2,"style":5,"label":"API","url":"https://openrouter.ai/meta-llama/llama-3.2-11b-vision-instruct/api"}]}]}}