{"data":{"capabilities":{"chat":true,"reasoning":true,"tools":true,"vision":false},"description":"Source-backed single-DGX-Spark K160 REAP profile with a 200K context claim, FP8 KV, MTP2, and full/piecewise CUDA graphs. The archived source's optional eager branch is deliberately absent. CLI launch stays blocked because the validated image is a host-local build rather than a published immutable artifact.","engine":{"graph_mode":"full-and-piecewise","name":"vllm","version":"0.1.dev17016+g27fd665bd.d20260526"},"facts":{"metadata":{"provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"registry-enrichment","url":"https://github.com/0xSero/local-ai-registry"}]},"reason":"not-observed","state":"unknown"}},"hardware_count":1,"hardware_id":"dgx-spark-gb10-128gb","id":"deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1","launch":{"accelerator_backend":"nvidia","compose":{"env_file":"env.example","file":"compose.yml","project_name":"deepseek-v4-flash-spark-200k-tp1"},"container":{"captured_at":"2026-08-27T06:04:13.773Z","compose_file":"compose.yml","digest":null,"image":null,"reason":"compose-manifest-resolves-image-outside-normalized-record","runtime":"docker-compose","source":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"recipe-launch","url":"https://github.com/0xSero/local-ai-registry"}],"state":"indirect"},"container_port":8000,"environment":{"CUDA_VISIBLE_DEVICES":"0","KV_CACHE_DTYPE":"fp8","VLLM_ENABLE_DEEPSEEK_V4_SPARSE_MLA_WARMUP":"0","VLLM_TRITON_MLA_SPARSE":"1","VLLM_TRITON_MLA_SPARSE_ALLOW_CUDAGRAPH":"1"},"host_port":8000,"ipc":"host","kind":"docker-compose","network_mode":"host","shm_size":"64g"},"metadata":{},"model_instance_id":"0xsero-deepseek-v4-flash-180b--nvfp4-mxfp4-experts","provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"normalized-recipe","url":"https://github.com/0xSero/local-ai-registry"}]},"recipe_source":"0xsero","schema_version":"local-ai-registry/v1","serving":{"kv_cache_tokens":537516,"max_concurrency":1,"max_context_tokens":200000,"tensor_parallel":1},"speed_sweep_ids":["deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep"],"status":"candidate","huggingface":{"link_type":"repository","provenance":{"captured_at":"2026-08-27T06:04:13.773Z","sources":[{"captured_at":"2026-08-27T06:04:13.773Z","kind":"huggingface-api","url":"https://huggingface.co/api/models/0xSero/DeepSeek-V4-Flash-180B"}]},"reason":"hf-api-confirmed-public","repository":"0xSero/DeepSeek-V4-Flash-180B","status":"known","url":"https://huggingface.co/0xSero/DeepSeek-V4-Flash-180B"},"registry":{"launchable":false,"speed_evidence":{"available":true,"count":1,"speed_sweep_ids":["deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep"],"detail_urls":["/api/v1/speed-sweep/deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep"]},"runtime":"docker"},"relationships":{"hardware":{"api":"/api/v1/hardware/dgx-spark-gb10-128gb","href":"/hardware/dgx-spark-gb10-128gb","id":"dgx-spark-gb10-128gb","name":"NVIDIA DGX Spark GB10"},"model":{"api":"/api/v1/models/deepseek-v4-flash-spark-200k","href":"/models/deepseek-v4-flash-spark-200k","id":"deepseek-v4-flash-spark-200k","name":"DeepSeek-V4-Flash-Spark-200K"},"model_instance":{"api":"/api/v1/model-instances/0xsero-deepseek-v4-flash-180b--nvfp4-mxfp4-experts","href":"/model-instances/0xsero-deepseek-v4-flash-180b--nvfp4-mxfp4-experts","id":"0xsero-deepseek-v4-flash-180b--nvfp4-mxfp4-experts","name":"0xSero/DeepSeek-V4-Flash-180B"},"speed_sweep":[{"api":"/api/v1/speed-sweep/deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep","href":"/speed-sweep/deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep","id":"deepseek-v4-flash-spark-200k-nvfp4-mxfp4-dgxspark-vllm-tp1-sweep"}]}},"meta":{"source":"registry"}}