[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"branding":3,"analytics":7,"article-gpt-models-didnt-fix-gender-bias-they-disguised-it":10,"sections":35},{"siteName":4,"siteTagline":5,"publisherName":4,"contactEmail":6},"The Revision","Tech news, decoded.","editor@therevision.news",{"gaMeasurementId":8,"adsenseClientId":9},"G-ZW2MV82GYR","ca-pub-8533917693782264",{"article":11},{"id":12,"slug":13,"title":14,"dek":15,"body_md":16,"tags_json":17,"published_at":18,"created_at":19,"updated_at":20,"status":21,"review_note":22,"review_notes":23,"image_url":22,"persona_id":22,"persona_name":22,"section":24,"tags":25,"sources":30,"feedback":34,"feedback_at":22,"cost_usd":34,"total_tokens":34},6924,"gpt-models-didnt-fix-gender-bias-they-disguised-it","GPT Models Didn't Fix Gender Bias, They Disguised It","A study of 450,000 GPT completions finds gender bias didn't vanish across model generations, it just evaded toxicity classifiers.","A new study argues that GPT models never got less biased about gender, they just got better at hiding it from the classifiers that grade them.\n\nResearchers analyzed 450,000 gender-directed completions across 15 models, from GPT-2 through GPT-5, under three demographic conditions. The sexual-violence content common in GPT-2's output about women mostly disappeared by GPT-4. But men-directed completions gained new positive territory, like caregiving and emotional range, that women-directed completions did not receive. By GPT-5, one 1,997-document cluster reframed breast cancer as a men's rights debate, with no comparable cluster on the women's side. Three separate toxicity classifiers scored all of it as non-toxic.\n\nThe sharper finding is structural: topic diversity in women-directed completions fell 36% relative to men's right around the GPT-4 safety overhaul, and representational harm disparity actually correlated with more recent release dates, while toxicity scores moved the opposite direction. That gap is the point. A falling toxicity score doesn't mean bias went away, it can mean the bias just moved somewhere the metric isn't looking.\n\nEvery model card that leans on toxicity scores alone as proof of progress is telling you less than it appears to.","[\"ai bias\",\"llm safety\",\"gender bias\",\"gpt\"]","2026-09-18T04:00:00.000Z","2026-09-18T22:52:38.914Z","2026-09-18T22:52:50.848Z","published",null,[],"ai",[26,27,28,29],"ai bias","llm safety","gender bias","gpt",[31],{"name":32,"url":33},"arXiv cs.AI","https:\u002F\u002Farxiv.org\u002Fabs\u002F2609.20779",0,{"sections":36},[37,40,45,50,55,59,63,68,72,77,82,87,92,97],{"name":38,"slug":24,"count":39,"latest_published_at":18},"AI",4082,{"name":41,"slug":42,"count":43,"latest_published_at":44},"Security","security",662,"2026-09-18T10:35:28.000Z",{"name":46,"slug":47,"count":48,"latest_published_at":49},"Policy","policy",339,"2026-09-17T12:00:00.000Z",{"name":51,"slug":52,"count":53,"latest_published_at":54},"Deals","deals",179,"2026-06-29T20:02:07.000Z",{"name":56,"slug":57,"count":58,"latest_published_at":18},"Hardware","hardware",155,{"name":60,"slug":61,"count":62,"latest_published_at":18},"Science","science",125,{"name":64,"slug":65,"count":66,"latest_published_at":67},"Consumer Tech","consumer-tech",99,"2026-09-09T17:27:33.000Z",{"name":69,"slug":70,"count":71,"latest_published_at":18},"Dev Tools","dev-tools",78,{"name":73,"slug":74,"count":75,"latest_published_at":76},"Software","software",75,"2026-09-10T20:41:21.000Z",{"name":78,"slug":79,"count":80,"latest_published_at":81},"Startups","startups",55,"2026-09-09T23:14:29.000Z",{"name":83,"slug":84,"count":85,"latest_published_at":86},"Gaming","gaming",43,"2026-09-10T12:18:06.000Z",{"name":88,"slug":89,"count":90,"latest_published_at":91},"General","general",41,"2026-09-08T01:57:23.000Z",{"name":93,"slug":94,"count":95,"latest_published_at":96},"Reviews","reviews",20,"2026-06-24T12:00:01.000Z",{"name":98,"slug":99,"count":100,"latest_published_at":101},"How-To","how-to",6,"2026-06-16T09:00:00.000Z"]