<?xml version="1.0" encoding="UTF-8"?>
<rss xmlns:dc="http://purl.org/dc/elements/1.1/" version="2.0">
  <channel>
    <title>InfoQ - AI Architecture</title>
    <link>https://www.infoq.com</link>
    <description>InfoQ AI Architecture feed</description>
    <item>
      <title>Pods as Workers, Not Agents: Rethinking the Deployment Unit for AI Agents on Kubernetes</title>
      <link>https://www.infoq.com/news/2026/08/pod-deployment-unit-ai-agents/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</link>
      <description>&lt;img src="https://res.infoq.com/news/2026/08/pod-deployment-unit-ai-agents/en/headerimage/header-1785763169197.jpeg"/&gt;&lt;p&gt;Running AI agents on Kubernetes raises a key question: should each agent get its own Pod? The kagent project argues no—agents are bursty, short-lived, can spawn subagents, and may wait for human approval, making one Pod per agent wasteful. Agent-substrate adds a control plane to schedule logical “Actors” onto long-lived worker Pods.&lt;/p&gt; &lt;i&gt;By Mark Silvester&lt;/i&gt;</description>
      <category>Cloud Native Computing Foundation</category>
      <category>Cloud Native Architecture</category>
      <category>Agents</category>
      <category>AI Architecture</category>
      <category>Kubernetes</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>Development</category>
      <category>Architecture &amp; Design</category>
      <category>DevOps</category>
      <category>news</category>
      <pubDate>Thu, 06 Aug 2026 06:00:00 GMT</pubDate>
      <guid>https://www.infoq.com/news/2026/08/pod-deployment-unit-ai-agents/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</guid>
      <dc:creator>Mark Silvester</dc:creator>
      <dc:date>2026-08-06T06:00:00Z</dc:date>
      <dc:identifier>/news/2026/08/pod-deployment-unit-ai-agents/en</dc:identifier>
    </item>
    <item>
      <title>Microsoft Agent Framework Harness and Hosted Agents Reach General Availability</title>
      <link>https://www.infoq.com/news/2026/08/agent-framework-harness-ga/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</link>
      <description>&lt;img src="https://res.infoq.com/news/2026/08/agent-framework-harness-ga/en/headerimage/generatedHeaderImage-1785148565256.jpg"/&gt;&lt;p&gt;Microsoft's Agent Framework now ships a supported runtime. Build 2026 brought the Agent Harness, the GitHub Copilot and Claude Agent SDK connectors, and the orchestration patterns to stable release; the harness and Foundry Hosted Agents have since reached GA. The shift is from an SDK for building agents to a governed platform for running them.&lt;/p&gt; &lt;i&gt;By Steef-Jan Wiggers&lt;/i&gt;</description>
      <category>Azure</category>
      <category>Cloud</category>
      <category>Large language models</category>
      <category>AI Architecture</category>
      <category>Development</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>Architecture &amp; Design</category>
      <category>DevOps</category>
      <category>news</category>
      <pubDate>Mon, 03 Aug 2026 10:30:00 GMT</pubDate>
      <guid>https://www.infoq.com/news/2026/08/agent-framework-harness-ga/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</guid>
      <dc:creator>Steef-Jan Wiggers</dc:creator>
      <dc:date>2026-08-03T10:30:00Z</dc:date>
      <dc:identifier>/news/2026/08/agent-framework-harness-ga/en</dc:identifier>
    </item>
    <item>
      <title>Presentation: Architecting AI Systems for the Messy Reality of Enterprises: Why Agentic Compute is the Missing Layer</title>
      <link>https://www.infoq.com/presentations/agentic-compute/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</link>
      <description>&lt;img src="https://res.infoq.com/presentations/agentic-compute/en/mediumimage/ArunJoseph-medium-1785311900135.jpeg"/&gt;&lt;p&gt;Arun Joseph shares real-world insights on scaling enterprise agentic platforms like Deutsche Telekom’s LMOS. He discusses bridging organizational fault lines, replacing tool sprawl with core platform abstractions, and moving beyond basic chatbots to operational intelligence systems through ephemeral agents and an Agent Definition Language (ADL).&lt;/p&gt; &lt;i&gt;By Arun Joseph&lt;/i&gt;</description>
      <category>InfoQ Dev Summit Munich 2025</category>
      <category>Agents</category>
      <category>AI Architecture</category>
      <category>Transcripts</category>
      <category>Architecture &amp; Design</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>presentation</category>
      <pubDate>Mon, 03 Aug 2026 08:08:00 GMT</pubDate>
      <guid>https://www.infoq.com/presentations/agentic-compute/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</guid>
      <dc:creator>Arun Joseph</dc:creator>
      <dc:date>2026-08-03T08:08:00Z</dc:date>
      <dc:identifier>/presentations/agentic-compute/en</dc:identifier>
    </item>
    <item>
      <title>AWS Launches Amazon GuardDuty Investigation Agent to Automate Threat Triage</title>
      <link>https://www.infoq.com/news/2026/07/guardduty-investigation-agent/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</link>
      <description>&lt;img src="https://res.infoq.com/news/2026/07/guardduty-investigation-agent/en/headerimage/generatedHeaderImage-1784791165676.jpg"/&gt;&lt;p&gt;AWS released a public preview of the GuardDuty investigation agent, which correlates findings, 90-day activity logs, and resource topologies into structured reports with risk ratings, confidence scores, and MITRE ATT&amp;CK classification. It is reachable through the AWS MCP Server, so investigations can run from agentic tooling. Preview quotas cap usage at 10 investigations per account per day.&lt;/p&gt; &lt;i&gt;By Steef-Jan Wiggers&lt;/i&gt;</description>
      <category>AWS</category>
      <category>Information Security</category>
      <category>Cloud</category>
      <category>AI Architecture</category>
      <category>Development</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>Architecture &amp; Design</category>
      <category>DevOps</category>
      <category>news</category>
      <pubDate>Tue, 28 Jul 2026 07:14:00 GMT</pubDate>
      <guid>https://www.infoq.com/news/2026/07/guardduty-investigation-agent/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</guid>
      <dc:creator>Steef-Jan Wiggers</dc:creator>
      <dc:date>2026-07-28T07:14:00Z</dc:date>
      <dc:identifier>/news/2026/07/guardduty-investigation-agent/en</dc:identifier>
    </item>
    <item>
      <title>Netflix Details its In-House LLM Serving Platform with Triton and vLLM</title>
      <link>https://www.infoq.com/news/2026/07/netflix-llm-platform/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</link>
      <description>&lt;img src="https://www.infoq.com/styles/static/images/logo/logo_bigger.jpg"/&gt;&lt;p&gt;Netflix has described the production lessons behind bringing LLM inference into its internal serving platform, including the challenges of supporting different model sizes, hardware requirements, and rapidly evolving inference engines.&lt;/p&gt; &lt;i&gt;By Matt Foster&lt;/i&gt;</description>
      <category>Infrastructure</category>
      <category>Open Source</category>
      <category>Large language models</category>
      <category>Model Inference</category>
      <category>AI Architecture</category>
      <category>Netflix</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>news</category>
      <pubDate>Mon, 27 Jul 2026 07:33:00 GMT</pubDate>
      <guid>https://www.infoq.com/news/2026/07/netflix-llm-platform/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=AI+Architecture</guid>
      <dc:creator>Matt Foster</dc:creator>
      <dc:date>2026-07-27T07:33:00Z</dc:date>
      <dc:identifier>/news/2026/07/netflix-llm-platform/en</dc:identifier>
    </item>
  </channel>
</rss>
