{
  "_meta": {
    "source": "https://apipriceindex.com/api/baseten/inkling-small",
    "attribution": "apipriceindex.com",
    "license": "CC BY 4.0",
    "license_url": "https://creativecommons.org/licenses/by/4.0/",
    "note": "Free to reuse with attribution to apipriceindex.com.",
    "generated_at": "2026-09-10T07:19:31.476Z"
  },
  "id": "baseten:inkling-small",
  "name": "Inkling Small (BaseTen)",
  "base_model": "inkling-small",
  "provider": {
    "id": "baseten",
    "name": "BaseTen",
    "url": "https://www.baseten.co"
  },
  "category": "llm-chat",
  "capabilities": [
    "text",
    "vision",
    "function-calling"
  ],
  "context_window": 1048576,
  "max_output": 32768,
  "pricing": {
    "model_id": "baseten:inkling-small",
    "input": 0.5,
    "output": 1.2,
    "cached_input": 0.1,
    "base_input": 0.5,
    "base_output": 1.2,
    "discount": 0,
    "status": 0,
    "source_url": "https://openrouter.ai/thinkingmachines/inkling-small",
    "verified_at": "2026-09-10T06:00:02Z",
    "sources": [
      {
        "name": "OpenRouter",
        "input": 0.5,
        "output": 1.2,
        "url": "https://openrouter.ai/thinkingmachines/inkling-small",
        "verified_at": "2026-09-10T06:00:02Z"
      },
      {
        "name": "LiteLLM",
        "input": 0.5,
        "output": 1.2,
        "url": "https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json",
        "verified_at": "2026-09-07T05:15:01Z"
      },
      {
        "name": "models.dev",
        "input": 0.5,
        "output": 1.2,
        "url": "https://models.dev/",
        "verified_at": "2026-09-07T05:30:01Z"
      },
      {
        "name": "Together",
        "input": 0.5,
        "output": 1.2,
        "url": "https://www.together.ai/pricing",
        "verified_at": "2026-09-07T05:40:02Z"
      }
    ],
    "confidence": "corroborated",
    "delta_pct": 0
  },
  "price_history": [
    [
      "2026-09-10",
      0.5,
      1.2
    ]
  ],
  "latency": [],
  "license": null,
  "cutoff": null,
  "intelligence": null,
  "quality": null,
  "api": {
    "base_url": "https://inference.baseten.co/v1",
    "model": "thinkingmachines/inkling-small",
    "via": "first-party",
    "key_env": "BASETEN_API_KEY",
    "docs_url": null
  },
  "facts": [
    "Inkling Small (BaseTen) costs $0.50 input / $1.20 output per 1M tokens.",
    "Cached input cuts the input price by 80%.",
    "A chatbot workload (1000 in / 500 out tokens) costs $1.10 per 1,000 requests.",
    "Cross-checked against LiteLLM's public price index, which lists inkling-small at $0.50/$1.20 per 1M (within 0% of this endpoint).",
    "Handles a 1,048,576-token context window, up to 32,768 output tokens.",
    "3% pricier than the cheapest inkling-small endpoint ($1.65 on DeepInfra)."
  ]
}