<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>KDA Blog</title>
    <link>https://nvlabs.github.io/kda/blog/</link>
    <description>Results, failure modes, and lessons from building agents that write, verify, and tune GPU kernels.</description>
    <language>en</language>
    <atom:link href="https://nvlabs.github.io/kda/blog/rss.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>KDA²: Kernel Design Agents (KDA) optimize Kimi Delta Attention (KDA)</title>
      <link>https://nvlabs.github.io/kda/blog/2026-09-27-kda-for-kda/</link>
      <guid isPermaLink="true">https://nvlabs.github.io/kda/blog/2026-09-27-kda-for-kda/</guid>
      <pubDate>Sun, 27 Sep 2026 00:00:00 GMT</pubDate>
      <description>Our agents wrote Kimi Delta Attention kernels that run up to 2.96× faster than FlashKDA on B300 with a tenth of its state error. Here is how, and how the agents tried to cheat along the way.</description>
      <category>Results</category>
      <category>Kimi Delta Attention</category>
      <category>Reward hacking</category>
    </item>
  </channel>
</rss>
