[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"branding":3,"analytics":7,"article-new-technique-helps-reasoning-ai-forget-more-cleanly":10,"sections":36},{"siteName":4,"siteTagline":5,"publisherName":4,"contactEmail":6},"The Revision","Tech news, decoded.","editor@therevision.news",{"gaMeasurementId":8,"adsenseClientId":9},"G-ZW2MV82GYR","ca-pub-8533917693782264",{"article":11},{"id":12,"slug":13,"title":14,"dek":15,"body_md":16,"tags_json":17,"published_at":18,"created_at":19,"updated_at":20,"status":21,"review_note":22,"review_notes":23,"image_url":24,"persona_id":22,"persona_name":22,"section":25,"tags":26,"sources":31,"feedback":35,"feedback_at":22,"cost_usd":35,"total_tokens":35},7075,"new-technique-helps-reasoning-ai-forget-more-cleanly","New Technique Helps Reasoning AI Forget More Cleanly","GUARD retrains reasoning models to stop leaking secrets in their thinking steps, not just polish the final answer.","AI models that think out loud before answering turn out to be lousy at truly forgetting things.\n\nResearchers have proposed a new method, called GUARD, to fix that for large reasoning models, whose chain-of-thought output makes forgetting harder because existing unlearning techniques only suppress a target fact or nudge internal representations without specifying what replaces it in that trace. That gap lets models hallucinate substitute details, produce malformed reasoning, or repeat themselves once the original content is blocked. GUARD instead trains models toward what its authors call a natural forgetting trajectory, a coherent, non-disclosing chain-of-thought followed by a stable refusal, built by distilling guided behavior from a frozen version of the model using guidance tokens, and the team also introduced a new metric, the Natural Forgetting Reasoning Score, to judge whether replacement reasoning is fluent and structurally sound rather than just checking if the secret leaked.\n\nThis matters because chain-of-thought is exactly where forgetting tends to fail quietly. A model can produce a clean, safe-looking final answer while its intermediate reasoning still spells out the protected fact or the unsafe instructions it was supposed to drop. As reasoning models get used in settings involving deletion requests, privacy takedowns, or safety filtering, scrubbing only the visible answer is a thin fix if the scratch work underneath still leaks.\n\nThe tests here cover two widely used distilled reasoning models and benchmarks the authors built or adapted themselves, which is a reasonable proof of concept but not independent confirmation that CoT unlearning is a solved problem.","[\"ai-safety\",\"unlearning\",\"reasoning-models\",\"chain-of-thought\"]","2026-09-21T04:00:00.000Z","2026-09-21T05:06:16.512Z","2026-09-21T05:06:28.195Z","published",null,[],"https:\u002F\u002Fcdn.xyz.onl\u002Farticle-images\u002Fnew-technique-helps-reasoning-ai-forget-more-cleanly.webp","ai",[27,28,29,30],"ai-safety","unlearning","reasoning-models","chain-of-thought",[32],{"name":33,"url":34},"arXiv cs.AI","https:\u002F\u002Farxiv.org\u002Fabs\u002F2609.21677",0,{"sections":37},[38,41,45,50,55,60,65,70,75,80,85,90,95,100],{"name":39,"slug":25,"count":40,"latest_published_at":18},"AI",4158,{"name":42,"slug":43,"count":44,"latest_published_at":18},"Security","security",679,{"name":46,"slug":47,"count":48,"latest_published_at":49},"Policy","policy",350,"2026-09-20T20:32:43.000Z",{"name":51,"slug":52,"count":53,"latest_published_at":54},"Deals","deals",179,"2026-06-29T20:02:07.000Z",{"name":56,"slug":57,"count":58,"latest_published_at":59},"Hardware","hardware",156,"2026-09-19T11:00:00.000Z",{"name":61,"slug":62,"count":63,"latest_published_at":64},"Science","science",130,"2026-09-20T13:48:11.000Z",{"name":66,"slug":67,"count":68,"latest_published_at":69},"Consumer Tech","consumer-tech",99,"2026-09-09T17:27:33.000Z",{"name":71,"slug":72,"count":73,"latest_published_at":74},"Dev Tools","dev-tools",78,"2026-09-18T04:00:00.000Z",{"name":76,"slug":77,"count":78,"latest_published_at":79},"Software","software",75,"2026-09-10T20:41:21.000Z",{"name":81,"slug":82,"count":83,"latest_published_at":84},"Startups","startups",55,"2026-09-09T23:14:29.000Z",{"name":86,"slug":87,"count":88,"latest_published_at":89},"Gaming","gaming",43,"2026-09-10T12:18:06.000Z",{"name":91,"slug":92,"count":93,"latest_published_at":94},"General","general",42,"2026-09-18T22:35:10.000Z",{"name":96,"slug":97,"count":98,"latest_published_at":99},"Reviews","reviews",20,"2026-06-24T12:00:01.000Z",{"name":101,"slug":102,"count":103,"latest_published_at":104},"How-To","how-to",6,"2026-06-16T09:00:00.000Z"]