{
  "_meta": {
    "source": "https://apipriceindex.com/api/google/gemini-3.6-flash",
    "attribution": "apipriceindex.com",
    "license": "CC BY 4.0",
    "license_url": "https://creativecommons.org/licenses/by/4.0/",
    "note": "Free to reuse with attribution to apipriceindex.com.",
    "generated_at": "2026-09-05T08:15:05.854Z"
  },
  "id": "google:gemini-3.6-flash",
  "name": "Gemini 3.6 Flash",
  "base_model": "gemini-3.6-flash",
  "provider": {
    "id": "google",
    "name": "Google",
    "url": "https://ai.google.dev"
  },
  "category": "llm-chat",
  "capabilities": [
    "text",
    "vision",
    "function-calling",
    "json-mode"
  ],
  "context_window": 1048576,
  "max_output": 65536,
  "pricing": {
    "model_id": "google:gemini-3.6-flash",
    "input": 0.75,
    "output": 3.75,
    "cached_input": 0.075,
    "source_url": "https://ai.google.dev/gemini-api/docs/pricing",
    "verified_at": "2026-09-03",
    "sources": [
      {
        "name": "Provider page",
        "input": 0.75,
        "output": 3.75,
        "url": "https://ai.google.dev/gemini-api/docs/pricing",
        "verified_at": "2026-09-03"
      },
      {
        "name": "LiteLLM",
        "input": 0.75,
        "output": 3.75,
        "url": "https://github.com/BerriAI/litellm/blob/main/model_prices_and_context_window.json",
        "verified_at": "2026-09-03T20:47:46Z"
      },
      {
        "name": "models.dev",
        "input": 0.75,
        "output": 3.75,
        "url": "https://models.dev/",
        "verified_at": "2026-09-04T10:52:49Z"
      }
    ],
    "confidence": "corroborated",
    "delta_pct": 0
  },
  "price_history": [
    [
      "2026-09-03",
      0.75,
      3.75
    ]
  ],
  "latency": [
    {
      "model_id": "google:gemini-3.6-flash",
      "region": "eu-paris",
      "ttft_ms": {
        "p50": 15070.2,
        "p95": 26605.1
      },
      "total_ms": {
        "p50": 15071.7,
        "p95": 26606.7
      },
      "tokens_per_s": 7729,
      "samples": 2,
      "error_rate": 0.6,
      "token_source": "usage",
      "measured_at": "2026-09-04T18:45:40Z"
    }
  ],
  "license": {
    "license": "Proprietary",
    "tier": "proprietary",
    "open_weights": false,
    "source_url": "https://ai.google.dev/gemini-api/terms"
  },
  "cutoff": {
    "cutoff": "2026-03",
    "label": "March 2026",
    "source_url": "https://deepmind.google/models/model-cards/gemini-3-6-flash/"
  },
  "intelligence": {
    "score": 51.6,
    "model_name": "Gemini 3.6 Flash",
    "verified_at": "2026-09-04",
    "benchmark_label": "Artificial Analysis Intelligence Index",
    "source_url": "https://artificialanalysis.ai/"
  },
  "quality": null,
  "api": {
    "base_url": "https://generativelanguage.googleapis.com/v1beta/openai",
    "model": "gemini-3.6-flash",
    "via": "first-party",
    "key_env": "GOOGLE_API_KEY",
    "docs_url": null
  },
  "facts": [
    "Gemini 3.6 Flash costs $0.75 input / $3.75 output per 1M tokens.",
    "Cached input cuts the input price by 90%.",
    "A chatbot workload (1000 in / 500 out tokens) costs $2.62 per 1,000 requests.",
    "Cross-checked: LiteLLM's public price index independently lists Google's gemini-3.6-flash at $0.75/$3.75 per 1M — corroborating the price shown here (within 0%).",
    "Handles a 1,048,576-token context window, up to 65,536 output tokens.",
    "Proprietary weights (Proprietary) — available via API only.",
    "Training knowledge cutoff: March 2026.",
    "Scores 51.6 on the Artificial Analysis Intelligence Index for gemini-3.6-flash.",
    "From eu-paris, p50 time-to-first-token is 15070 ms (2 samples).",
    "Intelligence value: 51.6 index at $4.50 blended = 11.5 index points per dollar — ranking #223 of 373 priced & AA-rated endpoints in the index.",
    "Throughput: 7729 tokens/s measured from eu-paris at $4.50 blended = 1717.6 tokens/s per dollar — ranking #3 of 11 priced & speed-tested endpoints in the index."
  ]
}