{"data":{"capabilities":{"chat":true,"reasoning":true,"tools":true,"vision":false},"description":"Local AI PostgreSQL candidate reconstructed from two completed GLM-5.2 evaluations on eight B200 GPUs. Model and image are now immutably pinned, but the source did not preserve graph-mode, Docker IPC/shared-memory settings, an on-hardware completion artifact, or a speed sweep, so CLI launch remains blocked.","engine":{"graph_mode":"unknown","name":"vllm","version":"0.23.0; image build 91df0fad4dc98a67c7659d9dbd915245d5c43d96"},"facts":{"metadata":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"not-observed","state":"unknown"},"serving.kv_cache_tokens":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"kv-cache-capacity-not-published","state":"unknown"},"serving.max_concurrency":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"not-observed","state":"unknown"},"serving.max_context_tokens":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"not-observed","state":"unknown"},"speed_sweep_ids":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"not-observed","state":"unknown"}},"hardware_count":8,"hardware_id":"b200-180gb","id":"glm-5.2-nvfp4-b200-vllm-tp8-mtp5","launch":{"accelerator_backend":"nvidia","arguments":["--model","nvidia/GLM-5.2-NVFP4","--served-model-name","nvidia/GLM-5.2-NVFP4","--host","0.0.0.0","--port","8000","--tensor-parallel-size","8","--pipeline-parallel-size","1","--trust-remote-code","--enable-chunked-prefill","--enable-prefix-caching","--enable-auto-tool-choice","--enable-prompt-tokens-details","--enable-force-include-usage","--enable-request-id-headers","--enable-log-requests","--max-num-seqs","1","--gpu-memory-utilization","0.90","--block-size","64","--language-model-only","--enable-expert-parallel","--max-model-len","1048576","--max-num-batched-tokens","8192","--tool-call-parser","glm47","--reasoning-parser","glm45","--kv-cache-dtype","fp8_e4m3","--speculative-config","{\"method\":\"mtp\",\"num_speculative_tokens\":5,\"rejection_sample_method\":\"standard\"}"],"container":{"captured_at":"2026-08-27T06:04:13.773Z","compose_file":null,"digest":"sha256:f03040c06dd43c0b48d0b471ae67edcc0c8fe8e63d4f762489d7be0015f527b2","image":"ghcr.io/davidmcc73/vllm-openai@sha256:f03040c06dd43c0b48d0b471ae67edcc0c8fe8e63d4f762489d7be0015f527b2","reason":"image-reference-in-launch","runtime":"docker","source":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"recipe-launch","url":"https://github.com/0xSero/local-ai-registry"}],"state":"digest-pinned"},"container_port":8000,"environment":{"HF_HOME":"/root/.cache/huggingface"},"host_port":8000,"image":"ghcr.io/davidmcc73/vllm-openai@sha256:f03040c06dd43c0b48d0b471ae67edcc0c8fe8e63d4f762489d7be0015f527b2","ipc":"host","kind":"docker","mounts":[{"read_only":false,"source":"~/.cache/huggingface","target":"/root/.cache/huggingface"}],"network_mode":"bridge","shm_size":"64g"},"metadata":{},"model_instance_id":"nvidia-glm-5-2-nvfp4--nvfp4","provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"normalized-recipe","url":"https://github.com/0xSero/local-ai-registry"}]},"recipe_source":"0xsero","schema_version":"local-ai-registry/v1","serving":{"kv_cache_tokens":null,"max_concurrency":null,"max_context_tokens":null,"tensor_parallel":8},"speed_sweep_ids":[],"status":"candidate","huggingface":{"link_type":"repository","provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"huggingface-api","url":"https://huggingface.co/api/models/nvidia/GLM-5.2-NVFP4"}]},"reason":"hf-api-confirmed-public","repository":"nvidia/GLM-5.2-NVFP4","status":"known","url":"https://huggingface.co/nvidia/GLM-5.2-NVFP4"},"registry":{"launchable":false,"speed_evidence":{"available":false,"count":0,"speed_sweep_ids":[],"detail_urls":[]},"runtime":"docker"},"relationships":{"hardware":{"api":"/api/v1/hardware/b200-180gb","href":"/hardware/b200-180gb","id":"b200-180gb","name":"NVIDIA B200"},"model":{"api":"/api/v1/models/glm-5-2","href":"/models/glm-5-2","id":"glm-5-2","name":"GLM-5.2"},"model_instance":{"api":"/api/v1/model-instances/nvidia-glm-5-2-nvfp4--nvfp4","href":"/model-instances/nvidia-glm-5-2-nvfp4--nvfp4","id":"nvidia-glm-5-2-nvfp4--nvfp4","name":"nvidia/GLM-5.2-NVFP4"},"speed_sweep":[]}},"meta":{"source":"registry"}}