<?xml version="1.0" encoding="UTF-8" ?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>Goutam Adwant | Technical Blog</title>
    <link>https://www.goutamadwant.com/blog</link>
    <description>Deep technical articles on AI infrastructure, distributed systems, LLM routing, and open-source engineering.</description>
    <language>en</language>
    <lastBuildDate>Thu, 10 Sep 2026 06:54:10 GMT</lastBuildDate>
    <atom:link href="https://www.goutamadwant.com/rss.xml" rel="self" type="application/rss+xml" />
    
    <item>
      <title><![CDATA[kvfleet Internals: Building a KV-Cache-Aware Routing Control Plane for LLM Fleets]]></title>
      <link>https://www.goutamadwant.com/blog/introducing-kvfleet</link>
      <guid isPermaLink="true">https://www.goutamadwant.com/blog/introducing-kvfleet</guid>
      <pubDate>Sat, 21 Mar 2026 00:00:00 GMT</pubDate>
      <description><![CDATA[A deep technical walkthrough of kvfleet's routing engine, from prompt fingerprinting and KV-cache affinity to policy-aware model selection, fallback chains, and OpenAI-compatible gateway mode.]]></description>
      <category>kvfleet</category>
      <category>llm-routing</category>
      <category>kv-cache</category>
      <category>python</category>
      <category>distributed-systems</category>
      <category>ai-infrastructure</category>
    </item>
  </channel>
</rss>