{"slug":"x402-agentutility-ai-voice-424f1c","title":"Turn written text into spoken audio you can drop straight into a player or app","host":"x402.agentutility.ai","method":"POST","resource":"https://x402.agentutility.ai/voice","category":"image","description":"Turn written text into spoken audio you can drop straight into a player or app. Send text (up to 4,000 characters) with an optional voice, model, speed, and output format of mp3, wav, opus, aac, or flac, and get back a hosted audio_url plus file size and content type, no base64 handling required. Ru","price_listed":0.05,"price_asked":null,"state":"effects","state_label":"Not tested: has real-world effects","checks_7d":0,"answered_7d":0,"latency_ms_median":null,"reported_calls_30d":2,"reported_payers_30d":1,"networks":["eip155:8453","solana:5eykt4UsFv8P8NJdTREpY1vzqKqZKvdp"],"badge":"unverified","paid_checks_7d":0,"paid_ok_7d":0,"example_input":{"body":{"text":"Hello there.","voice":"af_sky"},"bodyType":"json","method":"POST","type":"http"},"output_schema":{"$schema":"https://json-schema.org/draft/2020-12/schema","properties":{"input":{"additionalProperties":false,"properties":{"body":{"properties":{"format":{"description":"Output audio format: mp3, wav, opus, aac, or flac. Optional, default mp3.","type":"string"},"model":{"description":"TTS model id. Optional, default 'tts-kokoro'.","type":"string"},"speed":{"description":"Playback speed multiplier. Optional, 0.25-4, default 1.","type":"number"},"text":{"description":"Text to synthesize into speech. Max 4000 chars.","type":"string"},"voice":{"description":"Voice id, e.g. 'af_sky'. Optional, default 'af_sky'.","type":"string"}},"required":["text"]},"bodyType":{"enum":["json","form-data","text"],"type":"string"},"method":{"enum":["POST"],"type":"string"},"type":{"const":"http","type":"string"}},"required":["type","method","bodyType","body"],"type":"object"},"output":{"properties":{"example":{"properties":{"audio_url":{"type":"string"}},"type":"object"},"type":{"type":"string"}},"required":["type"],"type":"object"}},"required":["input"],"type":"object"},"history":[],"description_full":"Turn written text into spoken audio you can drop straight into a player or app. Send text (up to 4,000 characters) with an optional voice, model, speed, and output format of mp3, wav, opus, aac, or flac, and get back a hosted audio_url plus file size and content type, no base64 handling required. Runs on Venice's text-to-speech models with 30+ selectable voices and adjustable playback speed from 0.25x to 4x. Use it as a text-to-speech API, voice synthesis tool, or narration generator for voiceovers, IVR prompts, or accessibility read-aloud features.","last_updated":"2026-09-02T22:54:06.238Z","schemes":["exact"]}