<?xml version="1.0" encoding="UTF-8"?>
<rss xmlns:dc="http://purl.org/dc/elements/1.1/" version="2.0">
  <channel>
    <title>InfoQ - Performance - Presentations</title>
    <link>https://www.infoq.com</link>
    <description>InfoQ Performance Presentations feed</description>
    <item>
      <title>Presentation: From S3 to GPU in One Copy: Rethinking Data Loading for ML Training</title>
      <link>https://www.infoq.com/presentations/vortex-columnar-file-format-gpu-streaming/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</link>
      <description>&lt;img src="https://res.infoq.com/presentations/vortex-columnar-file-format-gpu-streaming/en/mediumimage/onur-satici-medium-1787813397611.jpeg"/&gt;&lt;p&gt;Onur Satici explains how Vortex, an open-source columnar file format under the Linux Foundation, revolutionizes high-throughput data loading. He details how cascading lightweight encodings, layout-based segment pruning, and zero-copy memory pipelines eliminate CPU/NVMe bottlenecks to stream S3 data straight to GPUs at speeds up to 60 Gbps without requiring upfront data reprocessing.&lt;/p&gt; &lt;i&gt;By Onur Satici&lt;/i&gt;</description>
      <category>Machine Learning</category>
      <category>Performance</category>
      <category>Transcripts</category>
      <category>Columnar Databases</category>
      <category>Rust</category>
      <category>CUDA</category>
      <category>S3</category>
      <category>QCon London 2026</category>
      <category>Data Lake</category>
      <category>Streaming</category>
      <category>GPU</category>
      <category>Data Pipelines</category>
      <category>Architecture</category>
      <category>Architecture &amp; Design</category>
      <category>Development</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>presentation</category>
      <pubDate>Fri, 04 Sep 2026 11:00:00 GMT</pubDate>
      <guid>https://www.infoq.com/presentations/vortex-columnar-file-format-gpu-streaming/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</guid>
      <dc:creator>Onur Satici</dc:creator>
      <dc:date>2026-09-04T11:00:00Z</dc:date>
      <dc:identifier>/presentations/vortex-columnar-file-format-gpu-streaming/en</dc:identifier>
    </item>
    <item>
      <title>Presentation: Beyond Prompting: Context Engineering for Production-Grade AI</title>
      <link>https://www.infoq.com/presentations/context-engineering-redis-llm-architecture/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</link>
      <description>&lt;img src="https://res.infoq.com/presentations/context-engineering-redis-llm-architecture/en/mediumimage/ricardo-ferreira-medium-1787820318843.jpg"/&gt;&lt;p&gt;Ricardo Ferreira discusses moving beyond simple prompt engineering to build production-grade AI applications. He shares practical architectural strategies for integrating long-term and short-term memory using Redis, managing LLM token limits via summarization, mitigating context rot with reranking and semantic caching, and controlling exponential API costs under strict latency constraints.&lt;/p&gt; &lt;i&gt;By Ricardo Ferreira&lt;/i&gt;</description>
      <category>Agents</category>
      <category>QCon AI Boston 2026</category>
      <category>LangChain4j</category>
      <category>Redis</category>
      <category>Large language models</category>
      <category>Generative AI</category>
      <category>Transcripts</category>
      <category>Caching</category>
      <category>Retrieval-Augmented Generation</category>
      <category>Architecture</category>
      <category>Architecture &amp; Design</category>
      <category>Development</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>presentation</category>
      <pubDate>Wed, 02 Sep 2026 11:00:00 GMT</pubDate>
      <guid>https://www.infoq.com/presentations/context-engineering-redis-llm-architecture/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</guid>
      <dc:creator>Ricardo Ferreira</dc:creator>
      <dc:date>2026-09-02T11:00:00Z</dc:date>
      <dc:identifier>/presentations/context-engineering-redis-llm-architecture/en</dc:identifier>
    </item>
    <item>
      <title>Presentation: Can Claude Fix Itself? Using LLMs for Incident Response</title>
      <link>https://www.infoq.com/presentations/claude-sre-incidents/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</link>
      <description>&lt;img src="https://res.infoq.com/presentations/claude-sre-incidents/en/mediumimage/alex-palcuie-medium-1786539336984.jpg"/&gt;&lt;p&gt;Anthropic reliability engineer Alex Palcuie shares practical lessons on using LLMs for real-world incident response. He explains where AI acts as a superhuman for observing logs and traces, why it still struggles with causation versus correlation during root-cause analysis, and how engineering leaders can integrate AI into on-call workflows without eroding human expertise.&lt;/p&gt; &lt;i&gt;By Alex Palcuie&lt;/i&gt;</description>
      <category>Site Reliability Engineering</category>
      <category>Incident Response</category>
      <category>Large language models</category>
      <category>Claude</category>
      <category>On-call</category>
      <category>QCon London 2026</category>
      <category>Automation</category>
      <category>Observability</category>
      <category>Transcripts</category>
      <category>Artificial Intelligence</category>
      <category>Architecture &amp; Design</category>
      <category>DevOps</category>
      <category>Development</category>
      <category>AI, ML &amp; Data Engineering</category>
      <category>presentation</category>
      <pubDate>Wed, 26 Aug 2026 11:00:00 GMT</pubDate>
      <guid>https://www.infoq.com/presentations/claude-sre-incidents/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</guid>
      <dc:creator>Alex Palcuie</dc:creator>
      <dc:date>2026-08-26T11:00:00Z</dc:date>
      <dc:identifier>/presentations/claude-sre-incidents/en</dc:identifier>
    </item>
    <item>
      <title>Presentation: Continuous Delivery for Foundational Platforms</title>
      <link>https://www.infoq.com/presentations/cd-reliability-innovation/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</link>
      <description>&lt;img src="https://res.infoq.com/presentations/cd-reliability-innovation/en/mediumimage/ian-nowland-medium-1787217621961.jpg"/&gt;&lt;p&gt;Ian Nowland discusses why conventional CI/CD practices break down for stateful, core infrastructure. Drawing from his leadership at AWS and Datadog, he shares actionable techniques for safe progressive deployments, synthetic testing in production, and mitigating blast radius in complex software platforms.&lt;/p&gt; &lt;i&gt;By Ian Nowland&lt;/i&gt;</description>
      <category>QCon San Francisco 2025</category>
      <category>Scalability</category>
      <category>Transcripts</category>
      <category>Platform Engineering</category>
      <category>Continuous Delivery</category>
      <category>DevOps</category>
      <category>presentation</category>
      <pubDate>Tue, 25 Aug 2026 13:34:00 GMT</pubDate>
      <guid>https://www.infoq.com/presentations/cd-reliability-innovation/?utm_campaign=infoq_content&amp;utm_source=infoq&amp;utm_medium=feed&amp;utm_term=Performance-presentations</guid>
      <dc:creator>Ian Nowland</dc:creator>
      <dc:date>2026-08-25T13:34:00Z</dc:date>
      <dc:identifier>/presentations/cd-reliability-innovation/en</dc:identifier>
    </item>
  </channel>
</rss>
