base_model = "google/gemini-2.5-pro" base_model_omit = ["structured_output"] reasoning_options = [{ type = "budget_tokens", min = 128, max = 32_768 }] [cost] input = 1.25 output = 10 cache_read = 0.125 [[cost.tiers]] tier = { type = "context", size = 200_000 } input = 2.5 output = 15 cache_read = 0.25