
  <rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
    <channel>
      <title>KI im Mittelstand – News &amp; Technologie-Übersicht</title>
      <link>https://www.ki-mittelstand.eu/blog</link>
      <description>Faktenbasierte News, Technologien und Praxislösungen zu Künstlicher Intelligenz für mittelständische Unternehmen. Ohne Hype, mit klaren Quellen.</description>
      <language>de-de</language>
      <managingEditor>phillip.pham@pexon-consulting.de (KI Mittelstand Team)</managingEditor>
      <webMaster>phillip.pham@pexon-consulting.de (KI Mittelstand Team)</webMaster>
      <lastBuildDate>Sat, 18 Jul 2026 00:00:00 GMT</lastBuildDate>
      <atom:link href="https://www.ki-mittelstand.eu/tags/performance/feed.xml" rel="self" type="application/rss+xml"/>
      
  <item>
    <guid>https://www.ki-mittelstand.eu/blog/ki-server-benchmarking-performance-messen</guid>
    <title>KI-Server-Benchmarking: Leistung messen und vergleichen</title>
    <link>https://www.ki-mittelstand.eu/blog/ki-server-benchmarking-performance-messen</link>
    <description>KI-Server-Benchmarking: Leistungsmetriken für LLM-Server messen — Latenz, Durchsatz, Token-Rates und VRAM-Usage im praktischen Benchmark.</description>
    <pubDate>Sat, 18 Jul 2026 00:00:00 GMT</pubDate>
    <author>phillip.pham@pexon-consulting.de (KI Mittelstand Team)</author>
    <category>benchmarking</category><category>performance</category><category>ki-server</category><category>llm-benchmark</category>
  </item>

  <item>
    <guid>https://www.ki-mittelstand.eu/blog/nvidia-triton-inference-server-produktion-2025-gpu-server-pr</guid>
    <title>ONNX Export: 3-5x schnellere KI-Inferenz</title>
    <link>https://www.ki-mittelstand.eu/blog/nvidia-triton-inference-server-produktion-2025-gpu-server-pr</link>
    <description>PyTorch-Modelle als ONNX exportieren: 3-5x schnellere Inferenz, 60 % weniger GPU-Kosten. BERT antwortet in 12 ms statt 45 ms.</description>
    <pubDate>Mon, 09 Mar 2026 00:00:00 GMT</pubDate>
    <author>phillip.pham@pexon-consulting.de (KI Mittelstand Team)</author>
    <category>onnx</category><category>inferenz</category><category>modelloptimierung</category><category>deployment</category><category>performance</category><category>deutschland</category>
  </item>

    </channel>
  </rss>
