[{"data":1,"prerenderedAt":660},["ShallowReactive",2],{"site-nav-content":3,"hiring-banner-content":163,"blog:/blog/what-to-look-for-in-model-routing":175,"blog-index-copy":619,"blog:/blog/what-to-look-for-in-model-routing:surround":640,"site-cta-content":641},{"header":4,"productNav":9,"nav":42,"footer":61,"askAI":116,"id":147,"title":148,"archived":149,"authors":150,"badge":150,"body":151,"date":150,"department":150,"description":155,"extension":158,"eyebrow":150,"faqHeader":150,"faqs":150,"footerBand":150,"headline":150,"image":150,"industry":150,"jobType":150,"listed":149,"location":150,"navigation":115,"openRoles":150,"pageLayout":150,"path":159,"relatedHeading":150,"seo":160,"series":150,"sitemap":115,"status":150,"stem":161,"subhead":150,"tags":150,"video":150,"whyJoin":150,"workplaceType":150,"__hash__":162},{"productLabel":5,"loginLabel":6,"contactLabel":7,"contactSalesLabel":8},"Product","Log in","Contact","Get started for free",[10,14,18,22,26,30,34,38],{"label":11,"to":12,"description":13},"Overview","/overview","Seven layers. One closed loop.",{"label":15,"to":16,"description":17},"Conflux","/product/conflux","Where your team, workstreams, and agents meet.",{"label":19,"to":20,"description":21},"Agent Teams","/product/agent-teams","Specialist teams - governed from day one.",{"label":23,"to":24,"description":25},"Lifecycle Graph","/product/lifecycle-graph","Intelligence that compounds across every interaction.",{"label":27,"to":28,"description":29},"Company Wiki","/product/wiki","Playbooks and policies where expertise stays.",{"label":31,"to":32,"description":33},"Workstreams","/product/workstreams","From brief to signed-off deliverable on one canvas.",{"label":35,"to":36,"description":37},"Perception Console","/product/perception","Ask your whole business in plain English.",{"label":39,"to":40,"description":41},"Governance","/product/governance","Frontier AI you can actually sign off on.",[43,46,49,52,55,58],{"label":44,"to":45},"Models","/models",{"label":47,"to":48},"Pricing","/pricing",{"label":50,"to":51},"Integrations","/integrations",{"label":53,"to":54},"Security","/security",{"label":56,"to":57},"Partners","/partners",{"label":59,"to":60},"Insights","/blog",{"productHeading":5,"companyHeading":62,"legalHeading":63,"docsLabel":64,"docsUrl":65,"statementLines":66,"copyright":69,"companyLinks":70,"legalLinks":85,"socialLinks":95,"bottomLinks":105},"Company","Legal","Docs","https://docs.gonimbus.ai",[67,68],"Stop training someone else's model.","Control your AI.","© 2026 Nimbus Intelligence, Inc. All rights reserved.",[71,72,73,74,75,77,80,83],{"label":47,"to":48},{"label":50,"to":51},{"label":53,"to":54},{"label":59,"to":60},{"label":76,"to":57},"Partner Program",{"label":78,"to":79},"Careers","/careers",{"label":81,"to":82},"System status","/status",{"label":7,"to":84},"/contact",[86,89,92],{"label":87,"to":88},"Terms of Service","/terms",{"label":90,"to":91},"Privacy Policy","/privacy",{"label":93,"to":94},"Compliance","/compliance",[96,99,102],{"label":97,"href":98},"LinkedIn","https://www.linkedin.com/company/gonimbusai/",{"label":100,"href":101},"X","https://x.com/gonimbusai",{"label":103,"href":104},"Instagram","https://www.instagram.com/gonimbus_ai/",[106,108,110,111,112],{"label":107,"to":88},"Terms",{"label":109,"to":91},"Privacy",{"label":93,"to":94},{"label":81,"to":82},{"label":113,"to":114,"external":115},"LLMs.txt","/llms.txt",true,{"text":117,"prompt":118},"Ask AI about Nimbus",{"I'm researching enterprise intelligence platforms and want to know how Nimbus combines perception, collaboration, and autonomous agents to drive strategic decision-making":119,"platforms":121},{" Summarize the highlights from Nimbus's website":120},"https://gonimbus.ai",[122,127,132,137,142],{"name":123,"label":124,"icon":125,"hrefPrefix":126},"chatgpt","ChatGPT","simple-icons:openai","https://chatgpt.com/?prompt=",{"name":128,"label":129,"icon":130,"hrefPrefix":131},"perplexity","Perplexity","mdi:magnify","https://www.perplexity.ai/search/new?q=",{"name":133,"label":134,"icon":135,"hrefPrefix":136},"grok","Grok","simple-icons:x","https://x.com/i/grok?text=",{"name":138,"label":139,"icon":140,"hrefPrefix":141},"claude","Claude","simple-icons:anthropic","https://claude.ai/new?q=",{"name":143,"label":144,"icon":145,"hrefPrefix":146},"google-ai","Google AI","simple-icons:google","https://www.google.com/search?udm=50&aep=11&q=","content/shared/nav.md","Site navigation",false,null,{"type":152,"value":153,"toc":154},"minimark",[],{"title":155,"searchDepth":156,"depth":156,"links":157},"",2,[],"md","/shared/nav",{"title":148,"description":155},"shared/nav","p6IsjEfjcrwkjwvGsPcuEkpoJSeUn2vJXPTxLWQfw9M",{"enabled":149,"message":164,"linkLabel":78,"linkHref":79,"id":165,"title":166,"archived":149,"authors":150,"badge":150,"body":167,"date":150,"department":150,"description":155,"extension":158,"eyebrow":150,"faqHeader":150,"faqs":150,"footerBand":150,"headline":150,"image":150,"industry":150,"jobType":150,"listed":149,"location":150,"navigation":115,"openRoles":150,"pageLayout":150,"path":171,"relatedHeading":150,"seo":172,"series":150,"sitemap":115,"status":150,"stem":173,"subhead":150,"tags":150,"video":150,"whyJoin":150,"workplaceType":150,"__hash__":174},"We're hiring! Join the team building the Sentient Enterprise.","content/shared/hiring.md","Hiring banner",{"type":152,"value":168,"toc":169},[],{"title":155,"searchDepth":156,"depth":156,"links":170},[],"/shared/hiring",{"title":166,"description":155},"shared/hiring","-6bioYD7lKYokGUVU3ff4hHTvB-sDyOMuCptKHnojfk",{"id":176,"title":177,"archived":149,"authors":178,"badge":181,"body":183,"date":608,"department":150,"description":609,"extension":158,"eyebrow":150,"faqHeader":150,"faqs":150,"footerBand":150,"headline":150,"image":150,"industry":150,"jobType":150,"listed":149,"location":150,"navigation":115,"openRoles":150,"pageLayout":150,"path":610,"relatedHeading":150,"seo":611,"series":612,"sitemap":115,"status":150,"stem":613,"subhead":150,"tags":614,"video":150,"whyJoin":150,"workplaceType":150,"__hash__":618},"content/blog/what-to-look-for-in-model-routing.md","What to Look for in Model Routing",[179],{"name":180,"to":120},"Nimbus Research",{"label":182},"Evaluation",{"type":152,"value":184,"toc":590},[185,193,196,218,223,273,284,288,291,311,319,342,350,353,357,398,408,411,416,419,448,451,455,458,465,474,478,482,488,492,495,499,502,506,513,517,524,528,542,546],[186,187,188,192],"p",{},[189,190,191],"strong",{},"Model routing"," is the policy that maps a task class to a model class before inference runs. It is not a brand preference. It is not a dropdown labelled “best.”",[186,194,195],{},"Someone using the flagship model to label a ticket is how you pay frontier prices for work a compact model could have finished in a second. That is not a moral failing. It is a missing policy. The product either chooses before the call, or a person chooses in a menu, or the default is the largest model “for quality.” Only the first is routing.",[186,197,198,205,206,211,212,217],{},[199,200,204],"a",{"href":201,"rel":202},"https://openai.com/api/pricing/",[203],"nofollow","OpenAI’s API pricing"," and ",[199,207,210],{"href":208,"rel":209},"https://www.anthropic.com/pricing",[203],"Anthropic’s pricing"," make the same point in public: compact and frontier models are not the same invoice line. ",[199,213,216],{"href":214,"rel":215},"https://hai.stanford.edu/ai-index/2025-ai-index-report",[203],"Stanford HAI’s 2025 AI Index"," has tracked how fast inference cost and capability moved — which is exactly why “best model” is not a routing policy. Best for a memo is not best for a classify step. Best last quarter is not best this quarter.",[219,220,222],"h2",{"id":221},"words-youll-hear","Words you’ll hear",[224,225,226,233,239,245,256,267],"ul",{},[227,228,229,232],"li",{},[189,230,231],{},"Frontier / flagship model."," The strongest (and usually most expensive) model a lab currently sells. Reserved for synthesis, hard reasoning, and novel language. Not for labelling.",[227,234,235,238],{},[189,236,237],{},"Compact / small model."," Faster and cheaper. Often enough for extract, classify, and summarise. “Small” is a cost and latency class, not an insult.",[227,240,241,244],{},[189,242,243],{},"Task class."," The kind of step: classify, retrieve, forecast, synthesise. If the platform cannot name the class, it cannot route. It can only default.",[227,246,247,250,251,255],{},[189,248,249],{},"NTU."," A metered unit of useful work so you can quote and cap a loop. See ",[199,252,254],{"href":253},"what-is-ai-token-economics","What is AI token economics",". Seats hide routing. NTU makes it visible.",[227,257,258,261,262,266],{},[189,259,260],{},"Model-agnostic."," The platform can call more than one provider. That is a menu, not a policy, until it ",[263,264,265],"em",{},"chooses"," by task class. Extra logos with a hidden flagship default is lock-in with branding.",[227,268,269,272],{},[189,270,271],{},"Always-flagship."," Marketing for “we use the best model.” A classify job does not need a long-context reasoner. Quality theatre is a cost event.",[186,274,275,276,279,280,283],{},"Routing is also not ",[189,277,278],{},"fine-tuning"," (changing a model’s weights), and it is not ",[189,281,282],{},"orchestration"," (what steps exist). You can orchestrate a brilliant graph and still send every node to the flagship. You can fine-tune a compact model and still need a policy that sends classify there. Evaluate them separately.",[219,285,287],{"id":286},"why-you-should-care","Why you should care",[186,289,290],{},"Teams do not wake up and choose waste. They inherit a default.",[224,292,293,299,305],{},[227,294,295,298],{},[189,296,297],{},"Single-model shop."," Every label, every search, every memo calls the same flagship. Finance sees one invoice and cannot split labelling from reasoning. You cannot cap what you cannot see.",[227,300,301,304],{},[189,302,303],{},"User-picked dropdown."," Power users pick the most expensive option “to be safe.” New hires copy that habit. Routing is now a training problem. Training problems do not survive quarter-end.",[227,306,307,310],{},[189,308,309],{},"Always-flagship as quality theatre."," Best for whom? Best for a forecast interpolation is often a time-series path, not a frontier model inventing a number that looks fine in a short demo.",[186,312,313,318],{},[199,314,317],{"href":315,"rel":316},"https://www.mckinsey.com/capabilities/quantumblack/our-insights/the-state-of-ai",[203],"McKinsey’s 2025 State of AI survey"," keeps showing the operational gap: regular use, then a struggle to scale because cost and workflow were never treated as a system. Routing is that system for inference. Without it, scale is a token bill.",[186,320,321,326,327,330,331,205,336,341],{},[199,322,325],{"href":323,"rel":324},"https://www.gartner.com/en/articles/ai-governance-trism",[203],"Gartner’s AI TRiSM framing"," implies you can ",[263,328,329],{},"see"," which model ran, on which data, at what cost. A platform that cannot show that is not ready, regardless of its red-team slides. ",[199,332,335],{"href":333,"rel":334},"https://eur-lex.europa.eu/eli/reg/2022/2554/oj",[203],"DORA",[199,337,340],{"href":338,"rel":339},"https://eur-lex.europa.eu/eli/dir/2022/2555/oj",[203],"NIS2"," change the evidence question: you should understand ICT dependencies. “We are not sure which model ran last Tuesday” is a dependency you cannot explain.",[186,343,344,345,349],{},"Choosing a model is choosing a brain. Choosing tools is choosing hands. Decide them separately. A compact model that extracts a refund still cannot write it without a quoted named signer if that is the workstream policy. Routing does not replace ",[199,346,348],{"href":347},"what-is-ai-governance","governance",". Governance does not replace routing. You need both.",[186,351,352],{},"Red flags: “we support many providers” with no task-class map; seat pricing that includes unlimited flagship; users pick the model in production; classify and memo share a model id in the demo; no NTU quote before a run expands; fallback is “switch the dropdown”; graph does not record model id per step; forecasting done by an LLM in the demo on a short series.",[219,354,356],{"id":355},"what-to-look-for","What to look for",[224,358,359,365,371,382,388],{},[227,360,361,364],{},[189,362,363],{},"They can refuse the flagship."," Run a classify-only job and show the model id. If classify used the same model as the memo, routing is a slide. Refusal is the proof. Support for many models is not.",[227,366,367,370],{},[189,368,369],{},"Task-class map you can read."," Classify → compact. Retrieve → embeddings, not stuffing hundreds of tickets into a long window. Forecast → a time-series path, not a frontier model interpolating a spreadsheet. Reason → frontier. If they cannot name the classes, they cannot route them.",[227,372,373,376,377,381],{},[189,374,375],{},"NTU quotes before the run expands."," Operators see an estimate and can set a workstream cap. Seat licences hide routing. Unlimited flagship under a seat is always-flagship with a predictable opex line. See ",[199,378,380],{"href":379},"total-cost-of-ownership-for-enterprise-ai","Total cost of ownership for enterprise AI",".",[227,383,384,387],{},[189,385,386],{},"Fallback is a logged promotion",", not “users will switch the dropdown.” Compact models fail on novel schemas and policy-edge language. Temporarily raise that class, budget-aware, then revert. A promotion without a log is a silent cost change. A dropdown is a training problem.",[227,389,390,393,394,381],{},[189,391,392],{},"The graph records the model id per step."," Six months later you can answer “which model drafted this?” without grepping provider dashboards. That is audit as well as cost. See ",[199,395,397],{"href":396},"how-to-evaluate-ai-audit-and-observability","How to evaluate AI audit and observability",[186,399,400,401,404,405,407],{},"Ask for four artefacts from one workstream run: model id per step; NTU per task class; a classify job that did ",[263,402,403],{},"not"," use the flagship; a forecast that did ",[263,406,403],{}," use an LLM as the estimator. If the vendor can only show a chat transcript and a blended token total, routing is not in the product.",[186,409,410],{},"Why each artefact matters: model id is Measure in NIST language. NTU per class is how Finance splits labelling from reasoning. Classify-without-flagship is the refusal test. Forecast-without-LLM is whether they know the difference between narration and estimation. Demos are short series. Production is seasonality, holidays, and missing days.",[412,413,415],"h3",{"id":414},"what-a-live-demo-should-prove","What a live demo should prove",[186,417,418],{},"Do not accept a provider logo wall.",[420,421,422,430,433,436,439,442,445],"ol",{},[227,423,424,425,429],{},"Run one ",[199,426,428],{"href":427},"what-is-an-ai-workstream","workstream"," with extract, classify, retrieve, and a memo.",[227,431,432],{},"Show model id per step. Classify and extract are compact. The memo may be frontier.",[227,434,435],{},"Show NTU (or tokens) per task class, quoted before the run grows, with a cap on the workstream.",[227,437,438],{},"Force a compact failure on a novel schema. Show a logged promotion, then a revert — not a user switching a dropdown.",[227,440,441],{},"Show a forecast path that is not an LLM interpolating a sheet. Narration can still be frontier.",[227,443,444],{},"Query the graph: which model drafted this payload? Answer from the ledger, not from a provider console.",[227,446,447],{},"Confirm a named-role gate still applies regardless of which model drafted. Routing chooses the brain. Governance still decides the write.",[186,449,450],{},"If they pass by opening a playground and picking “best,” you evaluated a dropdown.",[219,452,454],{"id":453},"how-this-shows-up-in-nimbus","How this shows up in Nimbus",[186,456,457],{},"Nimbus treats routing as an operating decision tied to workstream steps: task type, sensitivity, and cost — not “best everywhere.” Compact models handle extract. Frontier models are reserved for synthesis. Spend is NTU-metered, quoted per workstream, visible per step.",[186,459,460,461,464],{},"Release gates apply regardless of which model drafted the payload. Connectors stay read-only by default. Routing decides ",[263,462,463],{},"which brain"," reads them. Governance still decides whether anything writes.",[186,466,467,468,470,471,473],{},"See ",[199,469,44],{"href":45},". For the unit of account, ",[199,472,254],{"href":253},". Score the four artefacts above. The product claim is the policy, not the catalogue.",[219,475,477],{"id":476},"questions-people-actually-ask","Questions people actually ask",[412,479,481],{"id":480},"is-model-agnostic-the-same-as-routing","Is “model-agnostic” the same as routing?",[186,483,484,485,487],{},"No. Model-agnostic means more than one provider. Routing means it ",[189,486,265],{}," by task class, with a default that is cheap where cheap is correct. A hidden always-flagship default is lock-in with extra logos. Ask what happens if the operator never touches a dropdown. If the answer is flagship, you have your policy.",[412,489,491],{"id":490},"should-operators-ever-pick-a-model","Should operators ever pick a model?",[186,493,494],{},"Rarely, and as an override. Production operators should brief outcomes. If quality depends on each user knowing which model is good at JSON, you have staffed a routing department by accident. Overrides should be logged, budget-aware, and exceptional. A dropdown on every run is how always-flagship returns through the side door.",[412,496,498],{"id":497},"why-not-put-forecasting-in-the-llm-if-the-numbers-look-fine-in-the-demo","Why not put forecasting in the LLM if the numbers look fine in the demo?",[186,500,501],{},"Demos are short series. Production is seasonality, holidays, and missing days. Keep narration on the frontier model and estimation on a time-series path. A fluent number is not a control. The warehouse or the statistical path already owns the number. RAG plus a frontier model is for policy language, not for revenue by region.",[412,503,505],{"id":504},"what-if-legal-requires-a-single-approved-model-vendor","What if legal requires a single approved model vendor?",[186,507,508,509,512],{},"Routing still applies ",[189,510,511],{},"inside"," that vendor’s catalogue: compact vs frontier vs embedding. Single-vendor is a contracting constraint, not an excuse to max tokens. Model-agnostic is nice. Task-class mapping inside one catalogue is the control. Do not skip routing because the RFP named one lab.",[412,514,516],{"id":515},"does-nis2-or-dora-change-the-routing-question","Does NIS2 or DORA change the routing question?",[186,518,519,520,523],{},"They change the ",[189,521,522],{},"evidence"," question. DORA and NIS2 expect you to understand ICT dependencies. “We are not sure which model ran last Tuesday” is a dependency you cannot explain. Record model id per step on the graph. That is enough to start. You do not need a new product category. You need Measure.",[219,525,527],{"id":526},"related-reading","Related reading",[186,529,530,534,535,539,540,381],{},[199,531,533],{"href":532},"how-to-evaluate-ai-workstream-platforms","How to evaluate AI workstream platforms",", ",[199,536,538],{"href":537},"how-to-evaluate-an-enterprise-ai-operating-system","How to evaluate an enterprise AI operating system",", and ",[199,541,380],{"href":379},[219,543,545],{"id":544},"sources","Sources",[224,547,548,554,560,566,572,578,584],{},[227,549,550],{},[199,551,553],{"href":201,"rel":552},[203],"OpenAI API pricing",[227,555,556],{},[199,557,559],{"href":208,"rel":558},[203],"Anthropic pricing",[227,561,562],{},[199,563,565],{"href":214,"rel":564},[203],"Stanford HAI, 2025 AI Index Report",[227,567,568],{},[199,569,571],{"href":315,"rel":570},[203],"McKinsey, The state of AI in 2025",[227,573,574],{},[199,575,577],{"href":323,"rel":576},[203],"Gartner, AI TRiSM / AI governance",[227,579,580],{},[199,581,583],{"href":333,"rel":582},[203],"DORA (Regulation 2022/2554)",[227,585,586],{},[199,587,589],{"href":338,"rel":588},[203],"NIS2 (Directive 2022/2555)",{"title":155,"searchDepth":156,"depth":156,"links":591},[592,593,594,598,599,606,607],{"id":221,"depth":156,"text":222},{"id":286,"depth":156,"text":287},{"id":355,"depth":156,"text":356,"children":595},[596],{"id":414,"depth":597,"text":415},3,{"id":453,"depth":156,"text":454},{"id":476,"depth":156,"text":477,"children":600},[601,602,603,604,605],{"id":480,"depth":597,"text":481},{"id":490,"depth":597,"text":491},{"id":497,"depth":597,"text":498},{"id":504,"depth":597,"text":505},{"id":515,"depth":597,"text":516},{"id":526,"depth":156,"text":527},{"id":544,"depth":156,"text":545},"2026-08-17","Model routing is a policy that uses a cheaper model for simple steps and a stronger model only when the task needs it — not a dropdown labelled “best.”","/blog/what-to-look-for-in-model-routing",{"title":177,"description":609},"evaluation","blog/what-to-look-for-in-model-routing",[612,615,616,617],"models","routing","cost","JKYnpKrXCP09Op6E05HbukMEWqsQxAJrrxccWzNIUxk",{"hero":620,"id":622,"title":623,"archived":149,"authors":150,"badge":150,"body":624,"date":150,"department":150,"description":628,"extension":158,"eyebrow":629,"faqHeader":150,"faqs":150,"footerBand":630,"headline":150,"image":150,"industry":150,"jobType":150,"listed":149,"location":150,"navigation":115,"openRoles":150,"pageLayout":150,"path":60,"relatedHeading":636,"seo":637,"series":150,"sitemap":115,"status":150,"stem":638,"subhead":150,"tags":150,"video":150,"whyJoin":150,"workplaceType":150,"__hash__":639},{"filename":621},"u2221455217_Flat_design_of_a_futuristic_minimalist_landscape__5d589295-cdea-4ea9-a262-be766881accf_1.png","content/blog/index.md","Exploring the future of intelligence.",{"type":152,"value":625,"toc":626},[],{"title":155,"searchDepth":156,"depth":156,"links":627},[],"Deep dives into pre-cognitive intelligence, sentient enterprises, and the evolving landscape of AI-driven business transformation.","Latest Research",{"headline":631,"description":632,"primaryLabel":633,"primaryTo":634,"secondaryLabel":635,"secondaryTo":12},"Stay at the frontier.","Subscribe for product updates and new insights.","Subscribe","/newsletter","Explore the platform","More research",{"title":623,"description":628},"blog/index","eK1RCXdDW8nfLSyKRXGB1mJm9FAmhAO6GWXwNKOMNVE",[150,150],{"fold":642,"id":646,"title":647,"archived":149,"authors":150,"badge":150,"body":648,"date":150,"department":150,"description":155,"extension":158,"eyebrow":150,"faqHeader":150,"faqs":150,"footerBand":652,"headline":150,"image":150,"industry":150,"jobType":150,"listed":149,"location":150,"navigation":115,"openRoles":150,"pageLayout":150,"path":656,"relatedHeading":150,"seo":657,"series":150,"sitemap":115,"status":150,"stem":658,"subhead":150,"tags":150,"video":150,"whyJoin":150,"workplaceType":150,"__hash__":659},{"headline":643,"description":644,"primaryLabel":8,"primaryTo":645,"secondaryLabel":635,"secondaryTo":12},"Run frontier AI your business actually owns.","Governed agent swarms, 2,000+ integrations, and a knowledge graph that stays inside your walls. Free 7-day trial.","/checkout","content/shared/cta.md","Site CTAs",{"type":152,"value":649,"toc":650},[],{"title":155,"searchDepth":156,"depth":156,"links":651},[],{"headline":653,"description":654,"primaryLabel":8,"primaryTo":645,"secondaryLabel":655,"secondaryTo":84},"See what governed AI looks like on your stack.","Connect your tools, run a workstream, and keep every decision on your ledger - free for 7 days.","Talk to our team","/shared/cta",{"title":647,"description":155},"shared/cta","YHK6Fb8AvCPR1zZq7R_xiXUG0hwhP5UxHA8Ix52JQp4",1787194067138]