<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>The Minutes of Brad Gerstner on ai inference</title>
    <link>https://minutesof.com/brad-gerstner/on/ai-inference/</link>
    <description>Everything Brad Gerstner has said on ai inference: 2 verbatim quotes between August 2025 and July 2026, each with a timestamp and a link to the recording…</description>
    <language>en</language>
    <lastBuildDate>Mon, 31 Aug 2026 13:09:03 +0000</lastBuildDate>
    <atom:link href="https://minutesof.com/brad-gerstner/on/ai-inference/feed.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>Gerstner argues the difference between $3 and $15 inference costs is irrelevant when replacing a $200 per hour</title>
      <link>https://minutesof.com/q/9800ca2a-f280-488b-8c6c-73dd88f653f2/</link>
      <guid isPermaLink="true">https://minutesof.com/q/9800ca2a-f280-488b-8c6c-73dd88f653f2/</guid>
      <description>“The difference between spending $3 on a cheap model or $15 on an expensive model to replace a $200 an hour consultant. It&#x27;s just irrelevance. That inference cost difference is irrelevant” — All-In Podcast</description>
      <pubDate>Mon, 13 Jul 2026 20:30:36 +0000</pubDate>
      <category>ai inference</category>
    </item>
    <item>
      <title>Gerstner reports Google inference generation increased 100x in a year, from 9 trillion to 980 trillion tokens</title>
      <link>https://minutesof.com/q/dafd98a0-8753-40e3-ab50-2929f7b14312/</link>
      <guid isPermaLink="true">https://minutesof.com/q/dafd98a0-8753-40e3-ab50-2929f7b14312/</guid>
      <description>“A year ago, Google per month was doing about 9,000,000,000,000 tokens a month in terms of inference generation, right, compute generation. Today, it&#x27;s 980,000,000,000,000 tokens.” — CNBC Television</description>
      <pubDate>Thu, 28 Aug 2025 17:12:24 +0000</pubDate>
      <category>ai inference</category>
    </item>
  </channel>
</rss>
