{"request_id":"req_01M45X2TZW5FFGBDSCWX656ZWJ","technology":{"canonical_id":"tech_000000000000009YBMQKPCC32C","parent_canonical_id":"tech_000000000000009YBMQKJJYW7V","parent_id_zoho":"358446000000430331","name":"vLLM","slug":"vllm","requires_context":false,"popularity":null,"description":"Open-source engine for serving and running inference with large language models.","explanation":"vLLM is an inference and serving engine for language models. It provides an application-facing server and supports model execution in operator-managed environments. The record identifies the engine itself rather than a model served through it.","official_website_url":"https://vllm.ai/","official_documentation_url":"https://docs.vllm.ai/en/latest/","source_repository_url":"https://github.com/vllm-project/vllm","use_cases":[{"statement":"Running an operator-managed API server for supported models.","evidence_urls":["https://docs.vllm.ai/en/latest/getting_started/quickstart/"]},{"statement":"Serving language-model inference to applications.","evidence_urls":["https://docs.vllm.ai/en/latest/"]}],"strengths":[{"statement":"Designed for high-throughput and memory-efficient model serving.","evidence_urls":["https://github.com/vllm-project/vllm"]}],"limitations":[],"technology_license":{"name":"Apache License 2.0","scope":"vLLM repository code","spdx_id":"Apache-2.0","url":"https://github.com/vllm-project/vllm/blob/main/LICENSE","evidence_urls":["https://github.com/vllm-project/vllm/blob/main/LICENSE"]},"organizations":[{"name":"vLLM project","role":"maintainer","website_url":"https://vllm.ai/","evidence_urls":["https://github.com/vllm-project/vllm"]}],"certifications":null,"deployment_options":[{"type":"self_hosted","scope":"Operator-managed inference server","evidence_urls":["https://docs.vllm.ai/en/latest/getting_started/quickstart/"]}],"lifecycle":"active","revision":1,"updated_at":"2026-10-03T09:08:24.765195Z","reviewed_at":"2026-10-03T09:08:24.765195Z","source_freshness_at":"2026-10-03T00:00:00Z","aliases":null,"category_ids":["cat_2WB3ATK0PPYBKENHGJ83A1ZXQR"],"domain_ids":["dom_40T95R6SBK8B47KBASXES1YMWA"],"classification_kind":"runtime","sources":[{"url":"https://docs.vllm.ai/en/latest/","type":"official","title":"vLLM documentation","claim":"Documents inference and serving features.","retrieved_at":"2026-10-03T00:00:00Z"},{"url":"https://docs.vllm.ai/en/latest/getting_started/quickstart/","type":"official","title":"vLLM quickstart","claim":"Documents inference and API server setup.","retrieved_at":"2026-10-03T00:00:00Z"},{"url":"https://github.com/vllm-project/vllm","type":"official","title":"vLLM source repository","claim":"Official project repository describes the inference engine.","retrieved_at":"2026-10-03T00:00:00Z"},{"url":"https://github.com/vllm-project/vllm/blob/main/LICENSE","type":"official","title":"vLLM license","claim":"Repository code is licensed under Apache-2.0.","retrieved_at":"2026-10-03T00:00:00Z"},{"url":"https://vllm.ai/","type":"official","title":"vLLM","claim":"Official vLLM project home.","retrieved_at":"2026-10-03T00:00:00Z"}],"external_ids":[{"system":"beast","value":"358446000127741004"}],"provenance":"Published taxonomy release 2026.10.7"},"release_version":"2026.10.7"}
