<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>LLMPodium — LLM Leaderboard Updates</title>
    <link>https://llmpodium.com</link>
    <atom:link href="https://llmpodium.com/rss.xml" rel="self" type="application/rss+xml"/>
    <description>LLM rankings, benchmark updates and new model additions on LLMPodium — the definitive LLM leaderboard.</description>
    <language>en</language>
    <lastBuildDate>Mon, 24 Aug 2026 08:43:40 GMT</lastBuildDate>
    <item>
      <title>State of LLMs 2026: The Podium Report</title>
      <link>https://llmpodium.com/blog/state-of-llms-2026</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/state-of-llms-2026#2026-08-05</guid>
      <pubDate>Wed, 05 Aug 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Our annual data snapshot: who leads the Podium Score, how wide the open-weights gap is, where prices landed, and which benchmarks decided the year.</description>
    </item>
    <item>
      <title>Best LLM for Coding in 2026: SWE-Bench, LiveCodeBench and Terminal-Bench Leaders</title>
      <link>https://llmpodium.com/blog/best-llm-for-coding-2026</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/best-llm-for-coding-2026#2026-08-03</guid>
      <pubDate>Mon, 03 Aug 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Which AI models write the best code in 2026? We compare Claude Fable 5, Claude Mythos Preview, GPT-5.6 Sol and Kimi K3 across SWE-Bench Verified, LiveCodeBench and Terminal-Bench.</description>
    </item>
    <item>
      <title>Open-Weight vs Proprietary LLMs in 2026: How Big Is the Gap?</title>
      <link>https://llmpodium.com/blog/open-weight-vs-proprietary-2026</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/open-weight-vs-proprietary-2026#2026-08-02</guid>
      <pubDate>Sun, 02 Aug 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Kimi K3 scores 83.3 on the Podium Score — closer to frontier proprietary models than ever. We quantify the 2026 open-weights gap using aggregated leaderboard data.</description>
    </item>
    <item>
      <title>LLM Pricing Compared (2026): From $0.28 to $50 per Million Output Tokens</title>
      <link>https://llmpodium.com/blog/llm-pricing-compared-2026</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/llm-pricing-compared-2026#2026-08-01</guid>
      <pubDate>Sat, 01 Aug 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>A 178× spread: the cheapest and most expensive frontier LLMs in 2026, price-per-intelligence analysis, and where the value sweet spot is.</description>
    </item>
    <item>
      <title>The Fastest LLMs in 2026: Output Speed and Latency Compared</title>
      <link>https://llmpodium.com/blog/fastest-llms-2026</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/fastest-llms-2026#2026-07-30</guid>
      <pubDate>Thu, 30 Jul 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Gemini 3.5 Flash-Lite leads at 389 tokens/second. Full speed ranking of frontier models, and why speed is the most underrated benchmark.</description>
    </item>
    <item>
      <title>Added Qwen 3 235B</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-28</guid>
      <pubDate>Tue, 28 Jul 2026 09:00:00 GMT</pubDate>
      <category>model</category>
      <description>Added Alibaba Qwen 3 235B MoE model with hybrid thinking mode to all leaderboards.</description>
    </item>
    <item>
      <title>LiveCodeBench scores updated</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-25</guid>
      <pubDate>Sat, 25 Jul 2026 09:00:00 GMT</pubDate>
      <category>benchmark</category>
      <description>Updated LiveCodeBench scores for all models with latest contamination-free results.</description>
    </item>
    <item>
      <title>Added Llama 4 Maverick</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-20</guid>
      <pubDate>Mon, 20 Jul 2026 09:00:00 GMT</pubDate>
      <category>model</category>
      <description>Added Meta Llama 4 Maverick 400B MoE model with 1M context window support.</description>
    </item>
    <item>
      <title>Claude Opus 4 benchmark data</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-15</guid>
      <pubDate>Wed, 15 Jul 2026 09:00:00 GMT</pubDate>
      <category>update</category>
      <description>Added full benchmark suite for Claude Opus 4 including SWE-Bench and AIME scores.</description>
    </item>
    <item>
      <title>Multi-language support</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-10</guid>
      <pubDate>Fri, 10 Jul 2026 09:00:00 GMT</pubDate>
      <category>feature</category>
      <description>LLMPodium now available in 10 languages: EN, ZH, JA, KO, TH, RU, DE, ES, IT, FR.</description>
    </item>
    <item>
      <title>What Is an LLM Arena and Why Does It Matter?</title>
      <link>https://llmpodium.com/blog/arena-explained</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/arena-explained#2026-07-10</guid>
      <pubDate>Fri, 10 Jul 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Understanding the concept of LLM arenas — live, crowd-sourced evaluations that reveal which models people actually prefer.</description>
    </item>
    <item>
      <title>Added Gemini 2.5 Pro/Flash</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-05</guid>
      <pubDate>Sun, 05 Jul 2026 09:00:00 GMT</pubDate>
      <category>model</category>
      <description>Added Google Gemini 2.5 Pro and 2.5 Flash with thinking capabilities.</description>
    </item>
    <item>
      <title>Pricing data refresh</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-07-01</guid>
      <pubDate>Wed, 01 Jul 2026 09:00:00 GMT</pubDate>
      <category>update</category>
      <description>Updated pricing data for all models from official API documentation.</description>
    </item>
    <item>
      <title>LLM Benchmarking Best Practices in 2026</title>
      <link>https://llmpodium.com/blog/benchmarking-best-practices</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/benchmarking-best-practices#2026-07-01</guid>
      <pubDate>Wed, 01 Jul 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>How to avoid data contamination, account for prompt sensitivity, and build reliable evaluations.</description>
    </item>
    <item>
      <title>SWE-Bench Verified scores</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-06-28</guid>
      <pubDate>Sun, 28 Jun 2026 09:00:00 GMT</pubDate>
      <category>benchmark</category>
      <description>Updated SWE-Bench Verified scores for frontier models.</description>
    </item>
    <item>
      <title>Added DeepSeek R1</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-06-20</guid>
      <pubDate>Sat, 20 Jun 2026 09:00:00 GMT</pubDate>
      <category>model</category>
      <description>Added DeepSeek R1 reasoning model with chain-of-thought capabilities.</description>
    </item>
    <item>
      <title>Arena page launched</title>
      <link>https://llmpodium.com/changelog</link>
      <guid isPermaLink="false">https://llmpodium.com/changelog#2026-06-15</guid>
      <pubDate>Mon, 15 Jun 2026 09:00:00 GMT</pubDate>
      <category>feature</category>
      <description>New Arena page for side-by-side model comparison launched.</description>
    </item>
    <item>
      <title>Understanding LLM Evaluation Methodology</title>
      <link>https://llmpodium.com/blog/llm-evaluation-methodology</link>
      <guid isPermaLink="false">https://llmpodium.com/blog/llm-evaluation-methodology#2026-06-15</guid>
      <pubDate>Mon, 15 Jun 2026 09:00:00 GMT</pubDate>
      <category>Article</category>
      <description>Deep dive into how LLMPodium normalizes and weights scores across multiple public leaderboards.</description>
    </item>
  </channel>
</rss>
