{"slug":"the-stall-intuitek-ai-cap-vision-analyze-13229e","title":"Analyze any image URL using GPT-4o-mini vision","host":"the-stall.intuitek.ai","method":"GET","resource":"https://the-stall.intuitek.ai/cap/vision-analyze","category":"image","description":"Analyze any image URL using GPT-4o-mini vision. Returns structured analysis based on the mode: describe (full description), ocr (text extraction), chart (data/trend extraction), ui (interface analysis), identify (object/subject ID), or qa (answer a specific question about the image). Input must be a","price_listed":0.05,"price_asked":0.05,"state":"answering","state_label":"Answering","checks_7d":1,"answered_7d":1,"latency_ms_median":613,"reported_calls_30d":10,"reported_payers_30d":1,"networks":["eip155:8453"],"badge":"unverified","paid_checks_7d":0,"paid_ok_7d":0,"example_input":{"method":"GET","queryParams":{},"type":"http"},"output_schema":{"$schema":"https://json-schema.org/draft/2020-12/schema","properties":{"input":{"additionalProperties":false,"properties":{"method":{"enum":["GET"],"type":"string"},"queryParams":{"properties":{"detail":{"default":"auto","description":"OpenAI vision detail level. 'auto' (default): model decides based on image size. 'low': faster, cheaper, less detail (best for simple images). 'high': slower, more detail (best for charts, dense text, complex scenes).","enum":["low","high","auto"],"type":"string"},"mode":{"default":"describe","description":"Analysis mode: describe (full scene description), ocr (text extraction), chart (data/chart analysis), ui (UI screenshot analysis), identify (object/subject identification), qa (answer a specific question about the image — requires the 'question' parameter).","enum":["describe","ocr","chart","ui","identify","qa"],"type":"string"},"question":{"description":"For mode=qa only: the specific question to answer about the image. E.g., 'What is the total revenue shown in Q3?' or 'What does the error message say?'","maxLength":500,"type":"string"},"url":{"description":"Publicly accessible URL of the image to analyze. Must return image/jpeg, image/png, image/gif, or image/webp content-type. Max file size: 20MB.","type":"string"}},"required":[],"type":"object"},"type":{"const":"http","type":"string"}},"required":["type","method"],"type":"object"},"output":{"properties":{"example":{"properties":{"analysis":{"description":"AI-generated analysis of the image based on the selected mode.","type":"string"},"finish_reason":{"description":"OpenAI finish reason (stop, length, content_filter).","type":"string"},"mode":{"description":"The analysis mode that was applied.","type":"string"},"model":{"description":"OpenAI model used for analysis.","type":"string"},"tokens":{"properties":{"completion":{"type":["integer","null"]},"prompt":{"type":["integer","null"]},"total":{"type":["integer","null"]}},"type":"object"},"url":{"description":"The image URL that was analyzed.","type":"string"}},"type":"object"},"type":{"type":"string"}},"required":["type"],"type":"object"}},"required":["input"],"type":"object"},"history":[{"day":"2026-09-24","reachable":true,"status":402,"valid_402":true,"asked_usdc":0.05,"price_match":true,"latency_ms":613,"error":null}],"description_full":"Analyze any image URL using GPT-4o-mini vision. Returns structured analysis based on the mode: describe (full description), ocr (text extraction), chart (data/trend extraction), ui (interface analysis), identify (object/subject ID), or qa (answer a specific question about the image). Input must be a publicly accessible image URL (JPEG, PNG, GIF, WebP). $0.050/call.","last_updated":"2026-09-21T10:02:31.6Z","schemes":["exact"]}