{
  "api_url": "https://models.omg.bitop.dev/api/v1/models/llama-4-maverick.json",
  "deprecation_date": null,
  "description": "Open-weight multimodal mixture-of-experts model (17B active parameters, 128 experts).",
  "docs": "https://huggingface.co/meta-llama/Llama-4-Maverick-17B-128E-Instruct",
  "family": "llama-4",
  "home_provider": "openrouter",
  "id": "llama-4-maverick",
  "knowledge_cutoff": "2024-08-31",
  "license": "Llama 4 Community License",
  "limits": {
    "context": 1048576,
    "input": null,
    "output": null
  },
  "modalities": {
    "input": [
      "text",
      "image"
    ],
    "output": [
      "text"
    ]
  },
  "name": "Llama 4 Maverick",
  "offerings": [
    {
      "aliases": [],
      "current_price": {
        "batch": null,
        "batch_tiers": [],
        "currency": "USD",
        "effective_from": "2026-10-09",
        "not_applicable": [],
        "notes": "OpenRouter's listed price; its top endpoint serves 128K context and 16,384 output tokens.",
        "source": {
          "carried_over": [],
          "cites": null,
          "entry_key": "meta-llama/llama-4-maverick",
          "fetched_at": "2026-10-09T19:37:00Z",
          "kind": "openrouter",
          "sha256": "bdc084896dbbfa88715d7c2212f0239499fb25aca504b4e5e4612fa2558cbf68",
          "url": "https://openrouter.ai/api/v1/models",
          "version": null
        },
        "standard": {
          "cache_read_tokens": "0.05",
          "input_tokens": "0.1875",
          "output_tokens": "0.6525"
        },
        "standard_tiers": [],
        "units": {
          "cache_read_tokens": "per 1M tokens",
          "input_tokens": "per 1M tokens",
          "output_tokens": "per 1M tokens"
        }
      },
      "label": null,
      "limits": {
        "context": 128000,
        "input": null,
        "output": 16384
      },
      "price_history": [
        {
          "batch": null,
          "batch_tiers": [],
          "currency": "USD",
          "effective_from": "2026-10-09",
          "not_applicable": [],
          "notes": "OpenRouter's listed price; its top endpoint serves 128K context and 16,384 output tokens.",
          "source": {
            "carried_over": [],
            "cites": null,
            "entry_key": "meta-llama/llama-4-maverick",
            "fetched_at": "2026-10-09T19:37:00Z",
            "kind": "openrouter",
            "sha256": "bdc084896dbbfa88715d7c2212f0239499fb25aca504b4e5e4612fa2558cbf68",
            "url": "https://openrouter.ai/api/v1/models",
            "version": null
          },
          "standard": {
            "cache_read_tokens": "0.05",
            "input_tokens": "0.1875",
            "output_tokens": "0.6525"
          },
          "standard_tiers": [],
          "units": {
            "cache_read_tokens": "per 1M tokens",
            "input_tokens": "per 1M tokens",
            "output_tokens": "per 1M tokens"
          }
        }
      ],
      "provider": "openrouter",
      "provider_name": "OpenRouter",
      "upstream_id": "meta-llama/llama-4-maverick"
    }
  ],
  "open_weights": true,
  "reasoning": null,
  "release_date": null,
  "retirement_date": null,
  "status": "active",
  "tool_call": null,
  "url": "https://models.omg.bitop.dev/models/llama-4-maverick",
  "vendor": "Meta"
}
