<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>Machine Pidgin — News &amp; Field Notes</title>
    <link>https://machinepidgin.org/news</link>
    <description>Research updates, negative results, corrections, and open questions from Machine Pidgin.</description>
    <language>en-us</language>
    <atom:link href="https://machinepidgin.org/feed.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>A positive notation result disappeared under prompt audit</title>
      <link>https://machinepidgin.org/news/benchmark-002-notation-audit</link>
      <guid isPermaLink="true">https://machinepidgin.org/news/benchmark-002-notation-audit</guid>
      <pubDate>Mon, 03 Aug 2026 12:00:00 GMT</pubDate>
      <category>Benchmark 002 · audited result</category>
      <description>One unequal answer cue accounted for more than the preregistered gain. Across 19 equivalent tasks, vernacular scored 88.2% and formal notation 85.5%.</description>
    </item>
    <item>
      <title>SPEAR/0.2 improved exact adherence in a small synthetic pilot</title>
      <link>https://machinepidgin.org/news/benchmark-001-exact-adherence</link>
      <guid isPermaLink="true">https://machinepidgin.org/news/benchmark-001-exact-adherence</guid>
      <pubDate>Sun, 02 Aug 2026 12:00:00 GMT</pubDate>
      <category>Benchmark 001 · held-out pilot</category>
      <description>Across 16 held-out tasks and four model tiers, exact on-task adherence rose from 71.9% in prose to 89.1% with SPEAR/0.2, with important validity limits.</description>
    </item>
  </channel>
</rss>