[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"model-dinov2-small":3},{"id":4,"slug":5,"name":6,"hf_id":6,"organization":7,"architecture":8,"architecture_config":9,"parameters":13,"parameters_label":14,"context_length":15,"license":15,"downloads":16,"likes":17,"status":18,"description":15,"huggingface_url":19,"created_at":20,"seo":21,"category":15,"family":24,"score":26,"variants":27,"quantizations":28,"requirements":55,"gpu_recommendations":111,"fitting_gpus":174,"community":205},504,"dinov2-small","facebook\u002Fdinov2-small","facebook","dinov2",{"layers":10,"head_dim":11,"hidden_size":12},12,64,384,33521664,"33.5M",null,4404915,71,"published","https:\u002F\u002Fhuggingface.co\u002Ffacebook\u002Fdinov2-small","2026-08-13T18:12:39+00:00",{"title":6,"description":22,"image":15,"robots":23,"canonical_url":15},"AI 模型部署决策引擎——开源 AI 的 GPU 需求、云端价格与成本分析。","index, follow",{"name":25,"slug":7},"Facebook",0.0077,[],[29,34,39,44,48,52],{"format":30,"bits":31,"size_factor":32,"quality_loss":33},"AWQ",4,0.263,0.03,{"format":35,"bits":36,"size_factor":37,"quality_loss":38},"FP16",16,1,0,{"format":40,"bits":41,"size_factor":42,"quality_loss":43},"FP8",8,0.5,0.01,{"format":45,"bits":31,"size_factor":46,"quality_loss":47},"GGUF",0.25,0.05,{"format":49,"bits":31,"size_factor":50,"quality_loss":51},"GPTQ",0.275,0.04,{"format":53,"bits":41,"size_factor":42,"quality_loss":54},"INT8",0.02,{"FP16":56,"FP8":79,"INT8":89,"AWQ":93,"GPTQ":103,"GGUF":107},{"minimum":57,"production":67,"recommended":72},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":60,"kv_cache_gb":61,"vram_gb":62,"ram_gb":63,"disk_gb":64,"requirement_source":65,"confidence":66},"minimum",4096,0.06,0.09,2.46,32,39.1,"inferred",0.95,{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":60,"kv_cache_gb":10,"vram_gb":70,"ram_gb":63,"disk_gb":71,"requirement_source":65,"confidence":66},"production",16384,18.24,104.1,{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":60,"kv_cache_gb":76,"vram_gb":77,"ram_gb":63,"disk_gb":78,"requirement_source":65,"confidence":66},"recommended",8192,2,1.5,4.43,65.1,{"minimum":80,"production":83,"recommended":86},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":33,"kv_cache_gb":61,"vram_gb":81,"ram_gb":63,"disk_gb":82,"requirement_source":65,"confidence":66},2.42,39.05,{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":33,"kv_cache_gb":10,"vram_gb":84,"ram_gb":63,"disk_gb":85,"requirement_source":65,"confidence":66},18.19,104.05,{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":33,"kv_cache_gb":76,"vram_gb":87,"ram_gb":63,"disk_gb":88,"requirement_source":65,"confidence":66},4.39,65.05,{"minimum":90,"production":91,"recommended":92},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":33,"kv_cache_gb":61,"vram_gb":81,"ram_gb":63,"disk_gb":82,"requirement_source":65,"confidence":66},{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":33,"kv_cache_gb":10,"vram_gb":84,"ram_gb":63,"disk_gb":85,"requirement_source":65,"confidence":66},{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":33,"kv_cache_gb":76,"vram_gb":87,"ram_gb":63,"disk_gb":88,"requirement_source":65,"confidence":66},{"minimum":94,"production":97,"recommended":100},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":54,"kv_cache_gb":61,"vram_gb":95,"ram_gb":63,"disk_gb":96,"requirement_source":65,"confidence":66},2.4,39.03,{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":54,"kv_cache_gb":10,"vram_gb":98,"ram_gb":63,"disk_gb":99,"requirement_source":65,"confidence":66},18.17,104.03,{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":54,"kv_cache_gb":76,"vram_gb":101,"ram_gb":63,"disk_gb":102,"requirement_source":65,"confidence":66},4.37,65.03,{"minimum":104,"production":105,"recommended":106},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":54,"kv_cache_gb":61,"vram_gb":95,"ram_gb":63,"disk_gb":96,"requirement_source":65,"confidence":66},{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":54,"kv_cache_gb":10,"vram_gb":98,"ram_gb":63,"disk_gb":99,"requirement_source":65,"confidence":66},{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":54,"kv_cache_gb":76,"vram_gb":101,"ram_gb":63,"disk_gb":102,"requirement_source":65,"confidence":66},{"minimum":108,"production":109,"recommended":110},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":54,"kv_cache_gb":61,"vram_gb":95,"ram_gb":63,"disk_gb":96,"requirement_source":65,"confidence":66},{"scenario":68,"context_length":69,"batch_size":31,"weight_gb":54,"kv_cache_gb":10,"vram_gb":98,"ram_gb":63,"disk_gb":99,"requirement_source":65,"confidence":66},{"scenario":73,"context_length":74,"batch_size":75,"weight_gb":54,"kv_cache_gb":76,"vram_gb":101,"ram_gb":63,"disk_gb":102,"requirement_source":65,"confidence":66},{"recommended":112,"cheapest":130,"performance":139,"min_complexity":156,"cluster_options":170,"no_match":171,"required_vram_gb":101,"candidate_count":172,"usage":173},{"gpu_model":113,"gpu_model_id":114,"gpu_count":37,"total_vram_gb":36,"hour_price":115,"monthly_cost":116,"providers":117,"bandwidth_gbps":119,"fp16_tflops":120,"interconnect_score":121,"parallel_efficiency":37,"score":122,"tier":123,"perf_metric":124,"perf_label":125,"perf_value":126,"workload":15,"cost_metric":127},"RTX 5080",23,0.1339,97.75,[118],"vast",960,113,40,0.5677,"cheapest","tok\u002Fs","48000 tok\u002Fs",48000,{"effective_value":128,"effective_unit":124,"capacity_factor":129},33600,0.7,{"gpu_model":131,"gpu_model_id":132,"gpu_count":37,"total_vram_gb":36,"hour_price":133,"monthly_cost":134,"providers":135,"bandwidth_gbps":38,"fp16_tflops":38,"interconnect_score":121,"parallel_efficiency":37,"score":136,"tier":123,"perf_metric":124,"perf_label":137,"perf_value":38,"workload":15,"cost_metric":138},"Tesla V100",17,0.0289,21.1,[118],0.5624,"0 tok\u002Fs",{"effective_value":38,"effective_unit":124,"capacity_factor":129},{"gpu_model":140,"gpu_model_id":141,"gpu_count":41,"total_vram_gb":142,"hour_price":143,"monthly_cost":144,"providers":145,"bandwidth_gbps":147,"fp16_tflops":148,"interconnect_score":121,"parallel_efficiency":149,"score":150,"tier":151,"perf_metric":124,"perf_label":152,"perf_value":153,"workload":15,"cost_metric":154},"H200 SXM",11,1128,3.5,20440,[146],"lambda",4800,989,0.4,0.2197,"balanced","768000 tok\u002Fs",768000,{"effective_value":155,"effective_unit":124,"capacity_factor":129},537600,{"gpu_model":157,"gpu_model_id":158,"gpu_count":37,"total_vram_gb":159,"hour_price":160,"monthly_cost":161,"providers":162,"bandwidth_gbps":164,"fp16_tflops":17,"interconnect_score":121,"parallel_efficiency":37,"score":165,"tier":123,"perf_metric":124,"perf_label":166,"perf_value":167,"workload":15,"cost_metric":168},"RTX 3090",3,24,0.0622,45.41,[118,163],"tensordock",936,0.5443,"46800 tok\u002Fs",46800,{"effective_value":169,"effective_unit":124,"capacity_factor":129},32760,[],false,26,"value",[175,179,183,190,197],{"name":131,"vram_gb":36,"cheapest_hour":133,"cost":176,"provider":118,"estimated_tps":38},{"hourly":133,"daily":177,"monthly":134,"yearly":178},0.69,253.16,{"name":157,"vram_gb":159,"cheapest_hour":160,"cost":180,"provider":118,"estimated_tps":167},{"hourly":160,"daily":181,"monthly":161,"yearly":182},1.49,544.87,{"name":184,"vram_gb":36,"cheapest_hour":185,"cost":186,"provider":118,"estimated_tps":38},"RTX 4070S Ti",0.0678,{"hourly":185,"daily":187,"monthly":188,"yearly":189},1.63,49.49,593.93,{"name":191,"vram_gb":36,"cheapest_hour":192,"cost":193,"provider":118,"estimated_tps":38},"RTX 4080S",0.0685,{"hourly":192,"daily":194,"monthly":195,"yearly":196},1.64,50.01,600.06,{"name":198,"vram_gb":41,"cheapest_hour":199,"cost":200,"provider":118,"estimated_tps":204},"RTX 5060 Ti",0.0719,{"hourly":199,"daily":201,"monthly":202,"yearly":203},1.73,52.49,629.84,22400,{"posts_count":38,"benchmarks_count":38}]