<?xml version="1.0" encoding="UTF-8"?>
<rss xmlns:dc="http://purl.org/dc/elements/1.1/" version="2.0">
  <channel>
    <title>InfoQ - Model Evaluation - News</title>
    <link>https://www.infoq.com</link>
    <description>InfoQ Model Evaluation News feed</description>
    <item>
      <title>Anthropic's Claude Breaches Sandbox During Model Security Evaluations</title>
      <link>https://www.infoq.com/news/2026/08/claude-sandox-breach/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Model+Evaluation-news</link>
      <description>&lt;img src="https://www.infoq.com/styles/static/images/logo/logo_bigger.jpg"/&gt;&lt;p&gt;Anthropic conducted an audit of 141006 evaluation runs after OpenAI's sandbox escape disclosure. The review identified three incidents where Claude models accessed the internet due to misconfigurations. These incidents involved unauthorised attacks on live targets. Anthropic has suspended offensive evaluations and plans to enhance security measures and collaborate with external auditors.&lt;/p&gt; &lt;i&gt;By Olimpiu Pop&lt;/i&gt;</description>
      <category>AI Security</category>
      <category>Large language models</category>
      <category>Frontier Model</category>
      <category>Cloud Security</category>
      <category>Security Breach</category>
      <category>Model Evaluation</category>
      <category>DevOps</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>news</category>
      <pubDate>Thu, 13 Aug 2026 10:10:00 GMT</pubDate>
      <guid>https://www.infoq.com/news/2026/08/claude-sandox-breach/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Model+Evaluation-news</guid>
      <dc:creator>Olimpiu Pop</dc:creator>
      <dc:date>2026-08-13T10:10:00Z</dc:date>
      <dc:identifier>/news/2026/08/claude-sandox-breach/en</dc:identifier>
    </item>
  </channel>
</rss>
