{"slug":"flat-rate-llm-kikoribera03-workers-dev-v2-inferencia-896efd","title":"LLM inference at a flat $1.00 per call, with no account, no API key and no token","host":"flat-rate-llm.kikoribera03.workers.dev","method":"POST","resource":"https://flat-rate-llm.kikoribera03.workers.dev/v2/inferencia","category":"market","description":"LLM inference at a flat $1.00 per call, with no account, no API key and no token metering. POST a prompt, get the completion: the same price for 10 tokens or 4000, while metered gateways scale with usage. Automatic failover across several large models; typical response under 1s. Pay per call in USDC","price_listed":1.0,"price_asked":1.0,"state":"answering","state_label":"Answering","checks_7d":1,"answered_7d":1,"latency_ms_median":409,"reported_calls_30d":5,"reported_payers_30d":1,"networks":["eip155:137","eip155:42161","eip155:8453"],"badge":"unverified","paid_checks_7d":0,"paid_ok_7d":0,"example_input":{"body":{"prompt":"Summarise in one line what the x402 protocol is."},"bodyType":"json","method":"POST","type":"http"},"output_schema":{"$schema":"https://json-schema.org/draft/2020-12/schema","properties":{"input":{"additionalProperties":false,"properties":{"body":{"properties":{"maxTokens":{"description":"Output token cap, up to 4000. Does not change the price.","type":"number"},"prompt":{"description":"The request for the model. Up to 24000 characters.","type":"string"},"sistema":{"description":"Optional system instruction.","type":"string"}},"required":["prompt"]},"bodyType":{"enum":["json","form-data","text"],"type":"string"},"method":{"enum":["POST"],"type":"string"},"type":{"const":"http","type":"string"}},"required":["type","method","bodyType","body"],"type":"object"},"output":{"properties":{"example":{"type":"object"},"type":{"type":"string"}},"required":["type"],"type":"object"}},"required":["input"],"type":"object"},"history":[{"day":"2026-09-24","reachable":true,"status":402,"valid_402":true,"asked_usdc":1.0,"price_match":true,"latency_ms":409,"error":null}],"description_full":"LLM inference at a flat $1.00 per call, with no account, no API key and no token metering. POST a prompt, get the completion: the same price for 10 tokens or 4000, while metered gateways scale with usage. Automatic failover across several large models; typical response under 1s. Pay per call in USDC over x402 - nothing to sign up for.","last_updated":"2026-09-09T23:29:01.625Z","schemes":["exact"]}