<?xml version="1.0" encoding="utf-8"?>
<feed xmlns="http://www.w3.org/2005/Atom">
  <title>Art of Cyber AI — Research Feed</title>
  <subtitle>Research releases on reinforcement learning, LLM post-training, agent evaluation, and efficient inference.</subtitle>
  <link href="https://artofcyberai.com/feed.xml" rel="self" type="application/atom+xml"/>
  <link href="https://artofcyberai.com/" rel="alternate" type="text/html"/>
  <id>https://artofcyberai.com/</id>
  <updated>2026-08-24T12:00:00-07:00</updated>
  <author>
    <name>Vikram Kharvi</name>
    <uri>https://artofcyberai.com/</uri>
  </author>
  <entry>
    <title>Post-training Is a Stack</title>
    <link href="https://trainrl.com/post-training-stack" rel="alternate" type="text/html"/>
    <id>https://trainrl.com/post-training-stack</id>
    <published>2026-08-13T12:00:00-07:00</published>
    <updated>2026-08-24T12:00:00-07:00</updated>
    <summary>The full map of supervised tuning, synthetic data, preferences, reinforcement learning, safety, agent trajectories, distillation, and continual improvement.</summary>
  </entry>
</feed>
