[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"model-codegemma-2b":3},{"id":4,"slug":5,"name":6,"hf_id":7,"organization":8,"architecture":9,"architecture_config":10,"parameters":11,"parameters_label":12,"context_length":13,"license":14,"downloads":15,"likes":16,"status":17,"description":10,"huggingface_url":18,"created_at":19,"seo":20,"category":23,"family":26,"score":29,"variants":30,"quantizations":31,"requirements":58,"gpu_recommendations":137,"fitting_gpus":202,"community":238},72,"codegemma-2b","CodeGemma 2B","google\u002Fcodegemma-2b","google","transformer",null,2000000000,"2B",8192,"MIT",19078,100,"published","https:\u002F\u002Fhuggingface.co\u002Fgoogle\u002Fcodegemma-2b","2026-08-12T03:11:10+00:00",{"title":6,"description":21,"image":10,"robots":22,"canonical_url":10},"AI 模型部署决策引擎——开源 AI 的 GPU 需求、云端价格与成本分析。","index, follow",{"name":24,"slug":25},"Coding AI","coding",{"name":27,"slug":28},"Gemma","gemma",0.0011,[],[32,37,42,47,51,55],{"format":33,"bits":34,"size_factor":35,"quality_loss":36},"AWQ",4,0.263,0.03,{"format":38,"bits":39,"size_factor":40,"quality_loss":41},"FP16",16,1,0,{"format":43,"bits":44,"size_factor":45,"quality_loss":46},"FP8",8,0.5,0.01,{"format":48,"bits":34,"size_factor":49,"quality_loss":50},"GGUF",0.25,0.05,{"format":52,"bits":34,"size_factor":53,"quality_loss":54},"GPTQ",0.275,0.04,{"format":56,"bits":44,"size_factor":45,"quality_loss":57},"INT8",0.02,{"FP16":59,"FP8":82,"INT8":95,"AWQ":99,"GPTQ":112,"GGUF":125},{"minimum":60,"production":69,"recommended":76},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":63,"kv_cache_gb":45,"vram_gb":64,"ram_gb":65,"disk_gb":66,"requirement_source":67,"confidence":68},"minimum",4096,3.73,9.66,32,44.81,"estimated",0.95,{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":63,"kv_cache_gb":72,"vram_gb":73,"ram_gb":74,"disk_gb":75,"requirement_source":67,"confidence":68},"production",16384,64,122.74,184.11,109.81,{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":63,"kv_cache_gb":44,"vram_gb":79,"ram_gb":80,"disk_gb":81,"requirement_source":67,"confidence":68},"recommended",2,25.63,38.44,70.81,{"minimum":83,"production":87,"recommended":91},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":84,"kv_cache_gb":45,"vram_gb":85,"ram_gb":65,"disk_gb":86,"requirement_source":67,"confidence":68},1.86,7.31,41.91,{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":84,"kv_cache_gb":72,"vram_gb":88,"ram_gb":89,"disk_gb":90,"requirement_source":67,"confidence":68},120.17,180.26,106.91,{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":84,"kv_cache_gb":44,"vram_gb":92,"ram_gb":93,"disk_gb":94,"requirement_source":67,"confidence":68},23.16,34.75,67.91,{"minimum":96,"production":97,"recommended":98},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":84,"kv_cache_gb":45,"vram_gb":85,"ram_gb":65,"disk_gb":86,"requirement_source":67,"confidence":68},{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":84,"kv_cache_gb":72,"vram_gb":88,"ram_gb":89,"disk_gb":90,"requirement_source":67,"confidence":68},{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":84,"kv_cache_gb":44,"vram_gb":92,"ram_gb":93,"disk_gb":94,"requirement_source":67,"confidence":68},{"minimum":100,"production":104,"recommended":108},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":101,"kv_cache_gb":45,"vram_gb":102,"ram_gb":65,"disk_gb":103,"requirement_source":67,"confidence":68},0.97,6.17,40.51,{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":101,"kv_cache_gb":72,"vram_gb":105,"ram_gb":106,"disk_gb":107,"requirement_source":67,"confidence":68},118.93,178.4,105.51,{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":101,"kv_cache_gb":44,"vram_gb":109,"ram_gb":110,"disk_gb":111,"requirement_source":67,"confidence":68},21.98,32.97,66.51,{"minimum":113,"production":117,"recommended":121},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":114,"kv_cache_gb":45,"vram_gb":115,"ram_gb":65,"disk_gb":116,"requirement_source":67,"confidence":68},0.98,6.19,40.53,{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":114,"kv_cache_gb":72,"vram_gb":118,"ram_gb":119,"disk_gb":120,"requirement_source":67,"confidence":68},118.95,178.42,105.53,{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":114,"kv_cache_gb":44,"vram_gb":122,"ram_gb":123,"disk_gb":124,"requirement_source":67,"confidence":68},21.99,32.99,66.53,{"minimum":126,"production":129,"recommended":133},{"scenario":61,"context_length":62,"batch_size":40,"weight_gb":40,"kv_cache_gb":45,"vram_gb":127,"ram_gb":65,"disk_gb":128,"requirement_source":67,"confidence":68},6.22,40.56,{"scenario":70,"context_length":71,"batch_size":34,"weight_gb":40,"kv_cache_gb":72,"vram_gb":130,"ram_gb":131,"disk_gb":132,"requirement_source":67,"confidence":68},118.98,178.47,105.56,{"scenario":77,"context_length":13,"batch_size":78,"weight_gb":40,"kv_cache_gb":44,"vram_gb":134,"ram_gb":135,"disk_gb":136,"requirement_source":67,"confidence":68},22.02,33.04,66.56,{"recommended":138,"cheapest":156,"performance":166,"min_complexity":183,"cluster_options":198,"no_match":199,"required_vram_gb":109,"candidate_count":200,"usage":201},{"gpu_model":139,"gpu_model_id":78,"gpu_count":40,"total_vram_gb":65,"hour_price":140,"monthly_cost":141,"providers":142,"bandwidth_gbps":145,"fp16_tflops":146,"interconnect_score":147,"parallel_efficiency":40,"score":148,"tier":149,"perf_metric":150,"perf_label":151,"perf_value":152,"workload":10,"cost_metric":153},"RTX 5090",0.3222,235.21,[143,144],"vast","runpod",1792,209,40,0.6811,"cheapest","tok\u002Fs","1847.4 tok\u002Fs",1847.4,{"effective_value":154,"effective_unit":150,"capacity_factor":155},1293.2,0.7,{"gpu_model":157,"gpu_model_id":158,"gpu_count":78,"total_vram_gb":65,"hour_price":159,"monthly_cost":160,"providers":161,"bandwidth_gbps":41,"fp16_tflops":41,"interconnect_score":147,"parallel_efficiency":162,"score":163,"tier":149,"perf_metric":150,"perf_label":164,"perf_value":41,"workload":10,"cost_metric":165},"Tesla V100",17,0.0272,39.71,[143],0.8,0.6298,"0 tok\u002Fs",{"effective_value":41,"effective_unit":150,"capacity_factor":155},{"gpu_model":167,"gpu_model_id":168,"gpu_count":44,"total_vram_gb":169,"hour_price":170,"monthly_cost":171,"providers":172,"bandwidth_gbps":174,"fp16_tflops":175,"interconnect_score":147,"parallel_efficiency":176,"score":177,"tier":178,"perf_metric":150,"perf_label":179,"perf_value":180,"workload":10,"cost_metric":181},"H200 SXM",11,1128,3.5,20440,[173],"lambda",4800,989,0.4,0.2197,"balanced","15835.1 tok\u002Fs",15835.1,{"effective_value":182,"effective_unit":150,"capacity_factor":155},11084.6,{"gpu_model":184,"gpu_model_id":185,"gpu_count":40,"total_vram_gb":186,"hour_price":187,"monthly_cost":188,"providers":189,"bandwidth_gbps":191,"fp16_tflops":192,"interconnect_score":147,"parallel_efficiency":40,"score":193,"tier":149,"perf_metric":150,"perf_label":194,"perf_value":195,"workload":10,"cost_metric":196},"RTX 3090",3,24,0.0678,49.49,[143,190],"tensordock",936,71,0.6514,"964.9 tok\u002Fs",964.9,{"effective_value":197,"effective_unit":150,"capacity_factor":155},675.4,[],false,25,"value",[203,207,215,222,230],{"name":184,"vram_gb":186,"cheapest_hour":187,"cost":204,"provider":143,"estimated_tps":195},{"hourly":187,"daily":205,"monthly":188,"yearly":206},1.63,593.93,{"name":208,"vram_gb":186,"cheapest_hour":209,"cost":210,"provider":143,"estimated_tps":214},"RTX 4090",0.1344,{"hourly":209,"daily":211,"monthly":212,"yearly":213},3.23,98.11,1177.34,1039.2,{"name":216,"vram_gb":186,"cheapest_hour":217,"cost":218,"provider":143,"estimated_tps":214},"RTX 4090D",0.1602,{"hourly":217,"daily":219,"monthly":220,"yearly":221},3.84,116.95,1403.35,{"name":223,"vram_gb":186,"cheapest_hour":224,"cost":225,"provider":143,"estimated_tps":229},"L4",0.2015,{"hourly":224,"daily":226,"monthly":227,"yearly":228},4.84,147.1,1765.14,309.3,{"name":231,"vram_gb":186,"cheapest_hour":232,"cost":233,"provider":143,"estimated_tps":237},"A10",0.2409,{"hourly":232,"daily":234,"monthly":235,"yearly":236},5.78,175.86,2110.28,618.6,{"posts_count":41,"benchmarks_count":41}]