[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"branding":3,"analytics":7,"article-comed-teaches-ai-systems-when-to-ask-for-a-second-opinion":10,"sections":35},{"siteName":4,"siteTagline":5,"publisherName":4,"contactEmail":6},"The Revision","Tech news, decoded.","editor@therevision.news",{"gaMeasurementId":8,"adsenseClientId":9},"G-ZW2MV82GYR","ca-pub-8533917693782264",{"article":11},{"id":12,"slug":13,"title":14,"dek":15,"body_md":16,"tags_json":17,"published_at":18,"created_at":19,"updated_at":20,"status":21,"review_note":22,"review_notes":23,"image_url":22,"persona_id":22,"persona_name":22,"section":24,"tags":25,"sources":30,"feedback":34,"feedback_at":22,"cost_usd":34,"total_tokens":34},7595,"comed-teaches-ai-systems-when-to-ask-for-a-second-opinion","COMED Teaches AI Systems When to Ask for a Second Opinion","A new controller called COMED decides when an AI model needs a second opinion, boosting accuracy without asking every time.","A new academic system decides, one query at a time, whether an AI model should phone a friend before answering.\n\nResearchers behind COMED (Controlled Model Escalation for Multi-LLM Deliberation) built a controller that sits after a large language model gives its first answer. It checks the model's own consistency, how close the routing decision was, and a quick peer probe, then picks one of three paths: accept the answer, verify it, or escalate to a second model for full collaboration. That is a middle ground between two existing approaches - routers that pick one model and stop, and dense collaboration systems that consult every peer on every query regardless of need. Tested across 16 open-weight model setups on medical, scientific, and general reasoning benchmarks, COMED beat both approaches, with gains up to 10.7 percentage points on MedQA, while calling fewer models and using fewer tokens than dense collaboration. On frontier models, it lifted GPT-5.5's score on the HLE benchmark from 23.1% to 28.1%.\n\nThe interesting part is the paper's own admission that collaboration is not automatically good. Peer input can rescue a wrong answer, but it can just as easily talk a model out of a correct one - the researchers call this a rescue-harm trade-off. That framing matters more than the benchmark numbers: most multi-agent AI pitches assume more consultation is always better, when it can just as easily add noise and cost.\n\nStill, this is a benchmark result from open research, not a shipped product feature. Whether assistants or coding agents adopt this kind of selective escalation - instead of the current fashion for stacking more agents and more tokens at every query - remains to be seen.","[\"llm\",\"multi-model ai\",\"ai research\",\"inference\"]","2026-09-24T04:00:00.000Z","2026-09-24T08:42:13.146Z","2026-09-24T08:42:18.732Z","published",null,[],"ai",[26,27,28,29],"llm","multi-model ai","ai research","inference",[31],{"name":32,"url":33},"arXiv cs.AI","https:\u002F\u002Farxiv.org\u002Fabs\u002F2609.26913",0,{"sections":36},[37,40,44,49,54,59,64,69,74,79,84,89,94,99],{"name":38,"slug":24,"count":39,"latest_published_at":18},"AI",4424,{"name":41,"slug":42,"count":43,"latest_published_at":18},"Security","security",724,{"name":45,"slug":46,"count":47,"latest_published_at":48},"Policy","policy",380,"2026-09-23T22:53:43.000Z",{"name":50,"slug":51,"count":52,"latest_published_at":53},"Deals","deals",229,"2026-09-24T13:16:02.000Z",{"name":55,"slug":56,"count":57,"latest_published_at":58},"Hardware","hardware",175,"2026-09-24T12:10:00.000Z",{"name":60,"slug":61,"count":62,"latest_published_at":63},"Science","science",136,"2026-09-24T09:00:00.000Z",{"name":65,"slug":66,"count":67,"latest_published_at":68},"Consumer Tech","consumer-tech",119,"2026-09-24T13:31:43.000Z",{"name":70,"slug":71,"count":72,"latest_published_at":73},"Software","software",85,"2026-09-23T20:00:00.000Z",{"name":75,"slug":76,"count":77,"latest_published_at":78},"Dev Tools","dev-tools",79,"2026-09-22T22:21:13.000Z",{"name":80,"slug":81,"count":82,"latest_published_at":83},"Startups","startups",66,"2026-09-23T17:28:38.000Z",{"name":85,"slug":86,"count":87,"latest_published_at":88},"Gaming","gaming",45,"2026-09-22T15:35:06.000Z",{"name":90,"slug":91,"count":92,"latest_published_at":93},"General","general",43,"2026-09-21T23:48:56.000Z",{"name":95,"slug":96,"count":97,"latest_published_at":98},"Reviews","reviews",29,"2026-09-24T13:00:00.000Z",{"name":100,"slug":101,"count":102,"latest_published_at":103},"How-To","how-to",6,"2026-06-16T09:00:00.000Z"]