{
  "vertex_facts": [
    {
      "id": "vtx_context_window_parity",
      "model_id": null,
      "model_family": "all models",
      "category": "context window",
      "description": "Context window sizes for a given Claude model generation are the same across the direct API, Bedrock, and Vertex AI -- context length is a property of the model itself, not of the hosting cloud. Vertex-specific request/response payload limits imposed by Google Cloud's own API gateway are a separate, infrastructure-level constraint layered on top, distinct from the model's context window.",
      "source_url": "https://cloud.google.com/vertex-ai/generative-ai/docs/partner-models/use-claude",
      "created_at": "2026-07-02 08:29:41",
      "cite_as": "https://subagentvertex.com/api/vertex-facts/vtx_context_window_parity"
    }
  ]
}