freedom-intelligence/inventory/catalog/deepseek-ai__DeepSeek-R1-Distill-Qwen-32B__candidate.yaml

87 lines
3.1 KiB
YAML
Raw Normal View History

id: deepseek-ai__DeepSeek-R1-Distill-Qwen-32B__candidate
status: candidate
name: DeepSeek-R1-Distill-Qwen-32B
org: deepseek-ai
source:
kind: huggingface
url: https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Qwen-32B
revision: main
model_card_url: https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Qwen-32B
project_url: https://github.com/deepseek-ai/DeepSeek-R1
license:
spdx: MIT
url: https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Qwen-32B
allows_offline_retention: true
allows_local_ops: true
allows_fine_tune: true
notes: "R1 distill series MIT — confirm card at pin time."
profile:
summary: "Larger R1 distill for stronger local reason when 14B is not enough."
original_source: https://huggingface.co/deepseek-ai/DeepSeek-R1-Distill-Qwen-32B
use_cases:
- "Hard offline reasoning beyond 14B distill quality"
- "Local eval ceiling before considering full R1 MoE"
- "Heavier agent planning loops on multi-GPU / high-VRAM hosts"
- "FT experiments on reason-style traces at 32B"
sweet_spots:
- "Best dense open reason step between 14B distill and full R1"
- "Still denser/portable than 600B-class MoE"
- "MIT-friendly R1 lineage"
not_ideal_for:
- "Default always-on chat (too heavy)"
- "Hosts without ~20+ GB VRAM for Q4"
- "When 14B distill already saturates task quality"
capability_notes: >
W-tier optional. Approve only with VRAM + soft-quota headroom after R spine
and primary S pulls. Do not confuse with full DeepSeek-R1 MoE (separate S entry).
swot:
strengths:
- "Material reason quality jump over 14B distill"
- "Dense, so simpler serving than full MoE"
- "Clear provenance in R1 distill family"
weaknesses:
- "VRAM and latency cost; poor default chat model"
- "Still not full R1; diminishing returns vs S2 full weights"
opportunities:
- "Route only hardest local tasks here"
- "Bridge until multi-GPU can host full R1"
threats:
- "Quota competition with S1/S4 strategic weights"
- "Newer mid-large open reasoners may leapfrog"
size:
total_bytes: 0
total_human: "~65 GB fp16 / ~20 GB Q4 (estimate)"
hardware_class:
min_vram_gb_q4: 20
min_vram_gb_fp16: 64
notes: "T2T3; only if hardware envelope supports"
axes: [B, C]
priority: medium
collection:
approved_by: ""
approved_at: null
downloaded_at: null
downloaded_by: ""
storage_path: ""
brief_refs:
- research/2026-07-24-baseline-field-survey.md
reason: "P1/W stronger local reasoner — approve only with VRAM + quota headroom."
tags: [tier-w, reasoning, distill, deepseek]
companions:
- deepseek-ai__DeepSeek-R1-Distill-Qwen-14B__candidate
- deepseek-ai__DeepSeek-R1__strategic
notes: "Do not collect full DeepSeek-V3/R1 MoE under this id."
history:
- at: "2026-07-24"
event: nominated
by: baseline-survey
detail: "P1 recommendation from initial deep research."
- at: "2026-07-28"
event: profile_swot_added
by: grok
detail: "schema 0.2 profile + SWOT."
- at: "2026-07-28"
event: decision_pass
by: bernd
detail: "W — stay candidate; need VRAM+quota"