[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"branding":3,"analytics":7,"article-a-faster-way-to-measure-how-neural-networks-actually-learn":10,"sections":34},{"siteName":4,"siteTagline":5,"publisherName":4,"contactEmail":6},"The Revision","Tech news, decoded.","editor@therevision.news",{"gaMeasurementId":8,"adsenseClientId":9},"G-ZW2MV82GYR","ca-pub-8533917693782264",{"article":11},{"id":12,"slug":13,"title":14,"dek":15,"body_md":16,"tags_json":17,"published_at":18,"created_at":19,"updated_at":20,"status":21,"review_note":22,"review_notes":23,"image_url":22,"persona_id":22,"persona_name":22,"section":24,"tags":25,"sources":29,"feedback":33,"feedback_at":22,"cost_usd":33,"total_tokens":33},9266,"a-faster-way-to-measure-how-neural-networks-actually-learn","A Faster Way to Measure How Neural Networks Actually Learn","A new trick estimates neural tangent kernel statistics in a fraction of the time, making a once-impractical diagnostic usable at real model scale.","A new paper shows how to measure a neural network's learning geometry without reconstructing the giant matrix that describes it.\n\nThe Neural Tangent Kernel (NTK) is a mathematical object that captures how a network's parameters respond to training updates at a given point in time. Computing it directly for any real-sized model is usually impossible, since the memory and compute costs explode. The researchers instead use randomized trace estimation, a technique called Hutch++, to approximate useful NTK statistics, including trace, Frobenius norm, effective rank, and alignment, using only matrix-vector products from standard automatic differentiation. They tested the approach on multilayer perceptrons, recurrent GRUs, and a 410-million-parameter language transformer, reporting orders-of-magnitude speedups over direct computation.\n\nThat speed matters because NTK theory has mostly lived in papers about tiny toy networks. It has been too expensive to check whether its predictions hold up at the sizes people actually train. With these estimators, the team used NTK alignment to study rich versus lazy training in RNNs and as a regularizer during data-scarce knowledge distillation, where it modestly improved generalization.\n\nThe gains are real but modest, and that is the honest framing here. This is not a new training method, it is a cheaper ruler. If it holds up outside the three architectures tested, it could let researchers finally test NTK-based theories of generalization on the models people actually deploy, instead of toy networks and a wave of hands.","[\"ai\",\"neural-networks\",\"machine-learning-research\",\"neural-tangent-kernel\"]","2026-10-01T04:00:00.000Z","2026-10-02T05:36:09.320Z","2026-10-02T05:36:12.956Z","published",null,[],"ai",[24,26,27,28],"neural-networks","machine-learning-research","neural-tangent-kernel",[30],{"name":31,"url":32},"arXiv cs.AI","https:\u002F\u002Farxiv.org\u002Fabs\u002F2511.10796",0,{"sections":35},[36,39,43,47,52,57,61,66,71,75,80,85,90,95],{"name":37,"slug":24,"count":38,"latest_published_at":18},"AI",5671,{"name":40,"slug":41,"count":42,"latest_published_at":18},"Security","security",820,{"name":44,"slug":45,"count":46,"latest_published_at":18},"Policy","policy",430,{"name":48,"slug":49,"count":50,"latest_published_at":51},"Deals","deals",298,"2026-09-30T21:00:26.000Z",{"name":53,"slug":54,"count":55,"latest_published_at":56},"Hardware","hardware",196,"2026-09-30T13:00:00.000Z",{"name":58,"slug":59,"count":60,"latest_published_at":18},"Science","science",165,{"name":62,"slug":63,"count":64,"latest_published_at":65},"Consumer Tech","consumer-tech",149,"2026-09-30T22:57:11.000Z",{"name":67,"slug":68,"count":69,"latest_published_at":70},"Dev Tools","dev-tools",93,"2026-10-01T02:30:48.000Z",{"name":72,"slug":73,"count":69,"latest_published_at":74},"Software","software","2026-09-30T21:41:11.000Z",{"name":76,"slug":77,"count":78,"latest_published_at":79},"Startups","startups",84,"2026-09-30T20:39:09.000Z",{"name":81,"slug":82,"count":83,"latest_published_at":84},"Gaming","gaming",51,"2026-09-30T16:24:30.000Z",{"name":86,"slug":87,"count":88,"latest_published_at":89},"General","general",50,"2026-09-30T21:37:54.000Z",{"name":91,"slug":92,"count":93,"latest_published_at":94},"Reviews","reviews",31,"2026-09-28T14:31:34.000Z",{"name":96,"slug":97,"count":98,"latest_published_at":99},"How-To","how-to",6,"2026-06-16T09:00:00.000Z"]