<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom" xmlns:dc="http://purl.org/dc/elements/1.1/">
  <channel>
    <title>prokube.ai Blog</title>
    <link>https://prokube.ai/en/blog/</link>
    <description>Articles about sovereign AI infrastructure, Kubernetes, MLOps, and agentic workloads.</description>
    <language>en</language>
    <lastBuildDate>Wed, 19 Aug 2026 00:00:00 GMT</lastBuildDate>
    <atom:link href="https://prokube.ai/blog/rss.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>How to Run Qwen 3.8 27B on Kubernetes with KServe</title>
      <link>https://prokube.ai/en/blog/serving-qwen38-27b-kserve-vllm/</link>
      <guid isPermaLink="true">https://prokube.ai/en/blog/serving-qwen38-27b-kserve-vllm/</guid>
      <pubDate>Wed, 19 Aug 2026 00:00:00 GMT</pubDate>
      <dc:creator>Dr. Christian Geier</dc:creator>
      <description>A complete KServe setup for Qwen3.8 27B, including a reusable vLLM runtime, BF16 and FP8 deployment settings, H100 NVL validation, and the most important tuning knobs.</description>
      <category>Qwen3.8</category>
      <category>KServe</category>
      <category>vLLM</category>
      <category>Kubernetes</category>
      <category>LLMs</category>
    </item>
    <item>
      <title>AI Gateway, API Gateway, Gateway API, and friends: A Map Through the Gateway Confusion</title>
      <link>https://prokube.ai/en/blog/ai-gateway-api-gateway-gateway-api/</link>
      <guid isPermaLink="true">https://prokube.ai/en/blog/ai-gateway-api-gateway-gateway-api/</guid>
      <pubDate>Sun, 21 Jun 2026 00:00:00 GMT</pubDate>
      <dc:creator>Dr.-Ing. Henrik Steude</dc:creator>
      <description>Why &quot;gateway&quot; in cloud native and AI now means almost anything, and how to avoid losing the plot completely.</description>
      <category>AI Gateway</category>
      <category>Kubernetes</category>
      <category>Gateway API</category>
      <category>LLMs</category>
      <category>Agentic AI</category>
    </item>
    <item>
      <title>Self-Hosting Gemma 4 on Kubernetes with KServe and vLLM</title>
      <link>https://prokube.ai/en/blog/self-hosting-gemma-4-kubernetes-kserve-vllm/</link>
      <guid isPermaLink="true">https://prokube.ai/en/blog/self-hosting-gemma-4-kubernetes-kserve-vllm/</guid>
      <pubDate>Fri, 10 Apr 2026 00:00:00 GMT</pubDate>
      <dc:creator>Reyan Korel Erben</dc:creator>
      <dc:creator>Dr. Christian Geier</dc:creator>
      <description>How we ran the Gemma 4 31B instruction-tuned model on a single A100 80GB GPU using KServe and vLLM, including the custom runtime, PVC setup, tuning flags, verification, monitoring, and troubleshooting notes.</description>
      <category>Gemma 4</category>
      <category>KServe</category>
      <category>vLLM</category>
      <category>Kubernetes</category>
      <category>LLMs</category>
    </item>
  </channel>
</rss>
