<?xml version="1.0" encoding="UTF-8"?>
<resources count="5">
  <resource slug="tensorrt-llm" kind="ai-library"><name>TensorRT-LLM</name><url>https://github.com/NVIDIA/TensorRT-LLM</url><description>NVIDIA's optimized LLM inference.</description></resource>
  <resource slug="exllamav2" kind="ai-library"><name>ExLlamaV2</name><url>https://github.com/turboderp-org/exllamav2</url><description>Fast local inference kernels.</description></resource>
  <resource slug="modal" kind="build-platform"><name>Modal</name><url>https://modal.com</url><description>Serverless GPU/CPU compute.</description></resource>
  <resource slug="novita" kind="model-provider"><name>Novita AI</name><url>https://novita.ai</url><description>Serverless GPU model APIs.</description></resource>
  <resource slug="vllm" kind="model-runner" stars="92134"><name>vLLM</name><url>https://github.com/vllm-project/vllm</url><description>High-throughput GPU serving for open models.</description></resource>
</resources>