<?xml version="1.0" encoding="UTF-8"?>
<rss version="2.0">
  <channel>
    <title>Made with vLLM</title>
    <link>https://madewithwhat.net/vllm/</link>
    <description>A curated, daily-updated gallery of the best open-source projects built with vLLM, ranked by GitHub stars. Discover dashboards, UI kits, e-commerce, blogs and dev tools.</description>
    <language>en</language>
    <lastBuildDate>Wed, 22 Jul 2026 09:29:08 GMT</lastBuildDate>
    <generator>MadeWithWhat catalog</generator>
    <item>
      <title>RAG-LCC</title>
      <link>https://madewithwhat.net/vllm/project/rag-lcc/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/rag-lcc/</guid>
      <pubDate>Tue, 21 Jul 2026 11:09:09 GMT</pubDate>
      <description>A hands‑on RAG experimentation lab. Largely configurable with debug insights. Classification‑driven corpus construction, filter chains, document loading, chat interaction, Open WebUI integration. Experimental by design and not production‑ready. · 15 GitHub stars · by HarinezumIgel</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>tinycode</title>
      <link>https://madewithwhat.net/vllm/project/tinycode/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/tinycode/</guid>
      <pubDate>Tue, 21 Jul 2026 11:09:09 GMT</pubDate>
      <description>A slim, local-LLM-first AI coding assistant. TUI, Web UI, and desktop app. Runs air-gapped with Ollama, vLLM, or any OpenAI-compatible endpoint. · 10 GitHub stars · by bobbyjohnstx</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>dspark-vllm-gx10</title>
      <link>https://madewithwhat.net/vllm/project/dspark-vllm-gx10/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/dspark-vllm-gx10/</guid>
      <pubDate>Tue, 21 Jul 2026 11:09:09 GMT</pubDate>
      <description>Two-node DGX Spark/ASUS GX10 DeepSeek V4 Flash DSpark NVFP4 port for vLLM 0.25, with live dashboard and reproducible deployment. · 10 GitHub stars · by Anemll</description>
      <category>Dashboards</category>
    </item>
    <item>
      <title>FunASR</title>
      <link>https://madewithwhat.net/vllm/project/funasr/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/funasr/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Industrial-grade speech recognition toolkit: 170x realtime, 50+ languages, speaker diarization, emotion detection, streaming, and OpenAI-compatible API. · 19,219 GitHub stars · by modelscope</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>llama-cookbook</title>
      <link>https://madewithwhat.net/vllm/project/llama-cookbook/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/llama-cookbook/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Welcome to the Llama Cookbook! This is your go to guide for Building with Llama: Getting started with Inference, Fine-Tuning, RAG. We also show you how to solve end to end problems using Llama model family and using them on various provider services · 18,401 GitHub stars · by meta-llama</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>Halfrost-Field</title>
      <link>https://madewithwhat.net/vllm/project/halfrost-field/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/halfrost-field/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Source Code Deep Dives, System Design &amp; Engineering Blogs | Halfrost-Field ：、 · 13,206 GitHub stars · by halfrost</description>
      <category>Blogs</category>
    </item>
    <item>
      <title>AI-Research-SKILLs</title>
      <link>https://madewithwhat.net/vllm/project/ai-research-skills/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/ai-research-skills/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Comprehensive open-source library of AI research and engineering skills for any AI model. Package the skills and your claude code/codex/gemini agent will be an AI research agent with full horsepower. Maintained by Orchestra Research. · 10,689 GitHub stars · by Orchestra-Research</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>LMCache</title>
      <link>https://madewithwhat.net/vllm/project/lmcache/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/lmcache/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>LMCache: Supercharge Your LLM with the Fastest KV Cache Layer · 10,541 GitHub stars · by LMCache</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>OpenRLHF</title>
      <link>https://madewithwhat.net/vllm/project/openrlhf/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/openrlhf/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>An Easy-to-use, Scalable and High-performance Agentic RL Framework based on Ray (PPO &amp; DAPO &amp; REINFORCE++ &amp; VLM &amp; TIS &amp; vLLM &amp; Ray &amp; Async RL) · 9,785 GitHub stars · by OpenRLHF</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>inference</title>
      <link>https://madewithwhat.net/vllm/project/inference/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/inference/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Swap GPT for any LLM by changing a single line of code. Xinference lets you run open-source, speech, and multimodal models on cloud, on-prem, or your laptop — all through one unified, production-ready inference API. · 9,429 GitHub stars · by xorbitsai</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>dynamo</title>
      <link>https://madewithwhat.net/vllm/project/dynamo/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/dynamo/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A Datacenter Scale Distributed Inference Serving Framework · 7,479 GitHub stars · by ai-dynamo</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>Mooncake</title>
      <link>https://madewithwhat.net/vllm/project/mooncake/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/mooncake/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Mooncake is the serving platform for Kimi, a leading LLM service provided by Moonshot AI. · 5,816 GitHub stars · by kvcache-ai</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>kserve</title>
      <link>https://madewithwhat.net/vllm/project/kserve/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/kserve/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Standardized Distributed Generative and Predictive AI Inference Platform for Scalable, Multi-Framework Deployment on Kubernetes · 5,681 GitHub stars · by kserve</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>UltraRAG</title>
      <link>https://madewithwhat.net/vllm/project/ultrarag/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/ultrarag/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A Low-Code MCP Framework for Building Complex and Innovative RAG Pipelines · 5,643 GitHub stars · by OpenBMB</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>gpustack</title>
      <link>https://madewithwhat.net/vllm/project/gpustack/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/gpustack/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A GPU cluster manager for high-performance AI model serving (vLLM, SGLang) and on-demand SSH-accessible GPU instances. · 5,317 GitHub stars · by gpustack</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>sparrow</title>
      <link>https://madewithwhat.net/vllm/project/sparrow/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/sparrow/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Structured data extraction, instruction calling and agentic workflows with ML, LLM and Vision LLM · 5,181 GitHub stars · by katanaml</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>llama-swap</title>
      <link>https://madewithwhat.net/vllm/project/llama-swap/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/llama-swap/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Reliable model swapping for any local OpenAI/Anthropic compatible server - llama.cpp, vllm, etc · 4,992 GitHub stars · by mostlygeek</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>semantic-router</title>
      <link>https://madewithwhat.net/vllm/project/semantic-router/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/semantic-router/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Intelligent Mixture-of-Models Router for Efficient Heterogeneous LLMs Inference · 4,962 GitHub stars · by vllm-project</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>tiny-llm</title>
      <link>https://madewithwhat.net/vllm/project/tiny-llm/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/tiny-llm/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A course of learning LLM inference serving on Apple Silicon for systems engineers: build a tiny vLLM + Qwen. · 4,361 GitHub stars · by skyzh</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>LazyLLM</title>
      <link>https://madewithwhat.net/vllm/project/lazyllm/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/lazyllm/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Easiest and laziest way for building multi-agent LLMs applications. · 3,853 GitHub stars · by LazyAGI</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>FastDeploy</title>
      <link>https://madewithwhat.net/vllm/project/fastdeploy/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/fastdeploy/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>High-performance Inference and Deployment Toolkit for LLMs and VLMs based on PaddlePaddle · 3,702 GitHub stars · by PaddlePaddle</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>cascadeflow</title>
      <link>https://madewithwhat.net/vllm/project/cascadeflow/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/cascadeflow/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Cascading runtime for AI agents. Optimize cost, latency, quality, and policy decisions inside the agent loop. · 3,295 GitHub stars · by lemony-ai</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>Rapid-MLX</title>
      <link>https://madewithwhat.net/vllm/project/rapid-mlx/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/rapid-mlx/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>The fastest local AI engine for Apple Silicon. 4.2x faster than Ollama, 0.08s cached TTFT, 100% tool calling. 17 tool parsers, prompt cache, reasoning separation, cloud routing. Drop-in OpenAI replacement. Works with Claude Code, Cursor, Aider. · 3,262 GitHub stars · by raullenchai</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>ramalama</title>
      <link>https://madewithwhat.net/vllm/project/ramalama/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/ramalama/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>RamaLama is an open-source developer tool that simplifies the local serving of AI models from any source and facilitates their use for inference in production, all through the familiar language of containers. · 2,956 GitHub stars · by containers</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>vllm-ascend</title>
      <link>https://madewithwhat.net/vllm/project/vllm-ascend/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/vllm-ascend/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Community maintained hardware plugin for vLLM on Ascend · 2,496 GitHub stars · by vllm-project</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>auto-round</title>
      <link>https://madewithwhat.net/vllm/project/auto-round/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/auto-round/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A SOTA quantization algorithm for high-accuracy low-bit LLM inference, seamlessly optimized for CPU/XPU/CUDA, with multi-datatype support and full compatibility with vLLM, SGLang, and Transformers. · 1,519 GitHub stars · by intel</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>vllm-mlx</title>
      <link>https://madewithwhat.net/vllm/project/vllm-mlx/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/vllm-mlx/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>OpenAI and Anthropic compatible server for Apple Silicon. Run LLMs and vision-language models (Llama, Qwen-VL, LLaVA) with continuous batching, MCP tool calling, and multimodal support. Native MLX backend, 400+ tok/s. Works with Claude Code. · 1,428 GitHub stars · by waybarrios</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>local-studio</title>
      <link>https://madewithwhat.net/vllm/project/local-studio/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/local-studio/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Control panel for VLLM, Sglang, llama.cpp, exllamav3 · 1,359 GitHub stars · by sybil-solutions</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>InferenceX</title>
      <link>https://madewithwhat.net/vllm/project/inferencex/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/inferencex/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Open Source Continuous Inference Benchmark Research Platform — Kimi K2.7-Code, MiniMax M3, DeepSeekv4, GLM5 - GB200 NVL72 vs MI355X vs B200 vs GB300 NVL72 &amp; soon™ TPUv6e/v7/Trainium2/3 | — Kimi K2.7-Code、MiniMax M3、DeepSeekv4、GLM5 - GB200 NVL72 vs MI355X vs B200 vs GB300 NVL72，™ TPUv6e/v7/Trainium2/3 · 1,244 GitHub stars · by SemiAnalysisAI</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>kubeai</title>
      <link>https://madewithwhat.net/vllm/project/kubeai/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/kubeai/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>AI Inference Operator for Kubernetes. The easiest way to serve ML models in production. Supports VLMs, LLMs, embeddings, and speech-to-text. · 1,222 GitHub stars · by kubeai-project</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>BricksLLM</title>
      <link>https://madewithwhat.net/vllm/project/bricksllm/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/bricksllm/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Enterprise-grade API gateway that helps you monitor and impose cost or rate limits per API key. Get fine-grained access control and monitoring per user, application, or environment. Supports OpenAI, Azure OpenAI, Anthropic, vLLM, and open-source LLMs. · 1,217 GitHub stars · by bricks-cloud</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>GPTQModel</title>
      <link>https://madewithwhat.net/vllm/project/gptqmodel/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/gptqmodel/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>LLM model quantization (compression) toolkit with HW acceleration support for Nvidia, AMD, Intel GPU and Intel/AMD/Apple CPU via HF, vLLM, and SGLang. · 1,205 GitHub stars · by ModelCloud</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>prometheus-eval</title>
      <link>https://madewithwhat.net/vllm/project/prometheus-eval/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/prometheus-eval/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Evaluate your LLM&apos;s response with Prometheus and GPT4 · 1,102 GitHub stars · by prometheus-eval</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>kvcached</title>
      <link>https://madewithwhat.net/vllm/project/kvcached/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/kvcached/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Virtualized Elastic KV Cache for Dynamic GPU Sharing and Beyond · 1,097 GitHub stars · by ovg-project</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>tiny-vllm</title>
      <link>https://madewithwhat.net/vllm/project/tiny-vllm/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/tiny-vllm/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Build your own high performance LLM inference engine in C++ and CUDA - a smaller version of vLLM · 919 GitHub stars · by jmaczan</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>llm_note</title>
      <link>https://madewithwhat.net/vllm/project/llm-note/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/llm-note/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>LLM notes, including model inference, transformer model structure, and llm framework code analysis notes. · 883 GitHub stars · by harleyszhang</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>llmcord</title>
      <link>https://madewithwhat.net/vllm/project/llmcord/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/llmcord/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Make Discord your LLM frontend - Supports any OpenAI compatible API (OpenRouter, Ollama and more) · 814 GitHub stars · by jakobdylanc</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>UniRL</title>
      <link>https://madewithwhat.net/vllm/project/unirl/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/unirl/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>UniRL is a Framework for Unified Multimodal Model Reinforcement Learning · 810 GitHub stars · by Tencent-Hunyuan</description>
      <category>DevTools</category>
    </item>
    <item>
      <title>BambooAI</title>
      <link>https://madewithwhat.net/vllm/project/bambooai/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/bambooai/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>A Python library powered by Language Models (LLMs) for conversational data discovery and analysis. · 782 GitHub stars · by pgalko</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>local_ai_ocr</title>
      <link>https://madewithwhat.net/vllm/project/local-ai-ocr/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/local-ai-ocr/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>An local, offline (after initial setup), portable OCR software that can process images and PDF files, using DeepSeek-OCR-2 AI (running directly on your machine). · 773 GitHub stars · by th1nhhdk</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>LightCompress</title>
      <link>https://madewithwhat.net/vllm/project/lightcompress/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/lightcompress/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>[EMNLP 2024 &amp; AAAI 2026] A powerful toolkit for compressing large models including LLMs, VLMs, and video generative models. · 733 GitHub stars · by ModelTC</description>
      <category>DevTools</category>
    </item>
    <item>
      <title>vidur</title>
      <link>https://madewithwhat.net/vllm/project/vidur/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/vidur/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Accurate, large-scale, and extensible simulator for LLM inference Systems · 641 GitHub stars · by microsoft</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>efficientsam3</title>
      <link>https://madewithwhat.net/vllm/project/efficientsam3/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/efficientsam3/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>EfficientSAM3 compresses SAM3 into lightweight, edge-friendly models via progressive knowledge distillation for fast promptable concept segmentation and tracking. · 628 GitHub stars · by SimonZeng7108</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>FlashTTS</title>
      <link>https://madewithwhat.net/vllm/project/flashtts/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/flashtts/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>SparkTTS、OrpheusTTS，。 · 612 GitHub stars · by HuiResearch</description>
      <category>DevTools</category>
    </item>
    <item>
      <title>AI-Bank-Statement-Document-Automation-By-LLM-And-Personal-Finanical-Analysis-Prediction</title>
      <link>https://madewithwhat.net/vllm/project/ai-bank-statement-document-automation-by-llm-and-personal-finanical-analysis-prediction/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/ai-bank-statement-document-automation-by-llm-and-personal-finanical-analysis-prediction/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>AI Bank Statement Document Automation By LLM model and Personal Finanical Analysis · 597 GitHub stars · by johnsonhk88</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>RoboBrain</title>
      <link>https://madewithwhat.net/vllm/project/robobrain/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/robobrain/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>[CVPR 2025] RoboBrain: A Unified Brain Model for Robotic Manipulation from Abstract to Concrete. Official Repository. · 558 GitHub stars · by FlagOpen</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>verl-omni</title>
      <link>https://madewithwhat.net/vllm/project/verl-omni/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/verl-omni/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Multimodal RL training framework for diffusion &amp; omni models · 548 GitHub stars · by verl-project</description>
      <category>Docs</category>
    </item>
    <item>
      <title>crater</title>
      <link>https://madewithwhat.net/vllm/project/crater/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/crater/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Crater is a cloud-native AI training &amp; inference platform. · 542 GitHub stars · by raids-lab</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>openinfer</title>
      <link>https://madewithwhat.net/vllm/project/openinfer/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/openinfer/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>Pure Rust + CUDA LLM inference engine — no PyTorch, OpenAI-compatible, serves Qwen3 to Kimi-K2 · 535 GitHub stars · by openinfer-project</description>
      <category>AI &amp; ML</category>
    </item>
    <item>
      <title>restai</title>
      <link>https://madewithwhat.net/vllm/project/restai/</link>
      <guid isPermaLink="true">https://madewithwhat.net/vllm/project/restai/</guid>
      <pubDate>Wed, 15 Jul 2026 22:28:33 GMT</pubDate>
      <description>RESTai is an AIaaS (AI as a Service) open-source platform. Supports many public and local LLM suported by Ollama/vLLM/etc. Precise embeddings usage, tuning, analytics etc. Built-in image/audio generation with dynamic loading generators. Live chat deployment. Built-in block based graphical language. Prompt versioning and much more... · 512 GitHub stars · by apocas</description>
      <category>AI &amp; ML</category>
    </item>
  </channel>
</rss>