<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>Developers Digest - vLLM</title>
    <link>https://www.developersdigest.tech/tags/vllm</link>
    <description>2 items tagged vLLM on Developers Digest - blog posts, tools, guides, and tutorials.</description>
    <language>en</language>
    <lastBuildDate>Tue, 01 Sep 2026 12:50:04 GMT</lastBuildDate>
    <atom:link href="https://www.developersdigest.tech/tags/vllm/feed.xml" rel="self" type="application/rss+xml" />
    <image>
      <url>https://avatars.githubusercontent.com/u/124798203?v=4</url>
      <title>Developers Digest - vLLM</title>
      <link>https://www.developersdigest.tech/tags/vllm</link>
    </image>
    <item>
      <title><![CDATA[Ollama vs LM Studio vs vLLM vs llama.cpp: Picking a Local Runtime for Coding Agents]]></title>
      <link>https://www.developersdigest.tech/blog/local-llm-runtime-for-coding-agents-2026</link>
      <guid isPermaLink="true">https://www.developersdigest.tech/blog/local-llm-runtime-for-coding-agents-2026</guid>
      <description><![CDATA[A fair, sourced comparison of the four runtimes developers reach for when they want a coding agent talking to a model on their own hardware instead of an API: Ollama's convenience, LM Studio's GUI, vLLM's throughput, and llama.cpp's control. What each is actually for, and which to pick.]]></description>
      <pubDate>Thu, 09 Jul 2026 00:00:00 GMT</pubDate>
      <category>Local LLM</category>
      <category>Ollama</category>
      <category>vLLM</category>
      <category>llama.cpp</category>
      <category>AI Agents</category>
      <category>Coding Tools</category>
    </item>
    <item>
      <title><![CDATA[vLLM vs TGI vs SGLang: Which Inference Server to Self-Host]]></title>
      <link>https://www.developersdigest.tech/blog/vllm-vs-tgi-vs-sglang-inference-server-comparison</link>
      <guid isPermaLink="true">https://www.developersdigest.tech/blog/vllm-vs-tgi-vs-sglang-inference-server-comparison</guid>
      <description><![CDATA[A fair comparison of vLLM, TGI, SGLang, TensorRT-LLM, llama.cpp, and LMDeploy for self-hosted LLM inference - batching, quantization, hardware, and ops.]]></description>
      <pubDate>Thu, 09 Jul 2026 00:00:00 GMT</pubDate>
      <category>Inference</category>
      <category>vLLM</category>
      <category>Self-Hosting</category>
      <category>LLM Serving</category>
    </item>
  </channel>
</rss>