<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>The Minutes of Dylan Patel on cost reduction</title>
    <link>https://minutesof.com/dylan-patel/on/cost-reduction/</link>
    <description>Everything Dylan Patel has said on cost reduction: 2 verbatim quotes between March 2025 and July 2026, each with a timestamp and a link to the recording…</description>
    <language>en</language>
    <lastBuildDate>Mon, 31 Aug 2026 13:10:11 +0000</lastBuildDate>
    <atom:link href="https://minutesof.com/dylan-patel/on/cost-reduction/feed.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>Patel argues layering KV cache offload and multi-token prediction on open-source engines drastically cuts infe</title>
      <link>https://minutesof.com/q/6f64d719-d15d-49df-847e-a57b86cfaaa1/</link>
      <guid isPermaLink="true">https://minutesof.com/q/6f64d719-d15d-49df-847e-a57b86cfaaa1/</guid>
      <description>“So if you take the, you know, just open source inference engines off the shelf, that gets you a certain level of cost.” — RAISE Summit</description>
      <pubDate>Thu, 16 Jul 2026 16:44:57 +0000</pubDate>
      <category>cost reduction</category>
    </item>
    <item>
      <title>Patel states same-quality AI model output became 1200x cheaper than GPT-3 through successive model releases.</title>
      <link>https://minutesof.com/q/baf6aa96-f58f-4dc9-aa6e-f3c668d9d92e/</link>
      <guid isPermaLink="true">https://minutesof.com/q/baf6aa96-f58f-4dc9-aa6e-f3c668d9d92e/</guid>
      <description>“Anthropic released new models. Google released new models. Meta released new models. Right? And then now it is 1200x times it&#x27;s 1200x cheaper for the same output.” — MedBricks Webcast</description>
      <pubDate>Thu, 27 Mar 2025 13:55:08 +0000</pubDate>
      <category>cost reduction</category>
    </item>
  </channel>
</rss>
