[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"model-finbert":3},{"id":4,"slug":5,"name":6,"hf_id":6,"organization":7,"architecture":8,"architecture_config":9,"parameters":13,"parameters_label":14,"context_length":15,"license":16,"downloads":17,"likes":18,"status":19,"description":16,"huggingface_url":20,"created_at":21,"seo":22,"category":25,"family":28,"score":31,"variants":32,"quantizations":33,"requirements":60,"gpu_recommendations":118,"fitting_gpus":182,"community":213},500,"finbert","ProsusAI\u002Ffinbert","ProsusAI","bert",{"layers":10,"head_dim":11,"hidden_size":12},12,64,768,108375552,"108.4M",512,null,4486222,1220,"published","https:\u002F\u002Fhuggingface.co\u002FProsusAI\u002Ffinbert","2026-08-13T18:12:11+00:00",{"title":6,"description":23,"image":16,"robots":24,"canonical_url":16},"AI 模型部署决策引擎——开源 AI 的 GPU 需求、云端价格与成本分析。","index, follow",{"name":26,"slug":27},"LLM","llm",{"name":29,"slug":30},"Prosusai","prosusai",0.0212,[],[34,39,44,49,53,57],{"format":35,"bits":36,"size_factor":37,"quality_loss":38},"AWQ",4,0.263,0.03,{"format":40,"bits":41,"size_factor":42,"quality_loss":43},"FP16",16,1,0,{"format":45,"bits":46,"size_factor":47,"quality_loss":48},"FP8",8,0.5,0.01,{"format":50,"bits":36,"size_factor":51,"quality_loss":52},"GGUF",0.25,0.05,{"format":54,"bits":36,"size_factor":55,"quality_loss":56},"GPTQ",0.275,0.04,{"format":58,"bits":46,"size_factor":47,"quality_loss":59},"INT8",0.02,{"FP16":61,"FP8":84,"INT8":95,"AWQ":99,"GPTQ":109,"GGUF":113},{"minimum":62,"production":72,"recommended":77},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":65,"kv_cache_gb":66,"vram_gb":67,"ram_gb":68,"disk_gb":69,"requirement_source":70,"confidence":71},"minimum",4096,0.2,0.09,2.71,32,39.31,"inferred",0.95,{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":65,"kv_cache_gb":10,"vram_gb":75,"ram_gb":68,"disk_gb":76,"requirement_source":70,"confidence":71},"production",16384,19.78,104.31,{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":65,"kv_cache_gb":81,"vram_gb":82,"ram_gb":68,"disk_gb":83,"requirement_source":70,"confidence":71},"recommended",8192,2,1.5,4.94,65.31,{"minimum":85,"production":89,"recommended":92},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":86,"kv_cache_gb":66,"vram_gb":87,"ram_gb":68,"disk_gb":88,"requirement_source":70,"confidence":71},0.1,2.59,39.16,{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":86,"kv_cache_gb":10,"vram_gb":90,"ram_gb":68,"disk_gb":91,"requirement_source":70,"confidence":71},19.64,104.16,{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":86,"kv_cache_gb":81,"vram_gb":93,"ram_gb":68,"disk_gb":94,"requirement_source":70,"confidence":71},4.81,65.16,{"minimum":96,"production":97,"recommended":98},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":86,"kv_cache_gb":66,"vram_gb":87,"ram_gb":68,"disk_gb":88,"requirement_source":70,"confidence":71},{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":86,"kv_cache_gb":10,"vram_gb":90,"ram_gb":68,"disk_gb":91,"requirement_source":70,"confidence":71},{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":86,"kv_cache_gb":81,"vram_gb":93,"ram_gb":68,"disk_gb":94,"requirement_source":70,"confidence":71},{"minimum":100,"production":103,"recommended":106},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":52,"kv_cache_gb":66,"vram_gb":101,"ram_gb":68,"disk_gb":102,"requirement_source":70,"confidence":71},2.52,39.08,{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":52,"kv_cache_gb":10,"vram_gb":104,"ram_gb":68,"disk_gb":105,"requirement_source":70,"confidence":71},19.57,104.08,{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":52,"kv_cache_gb":81,"vram_gb":107,"ram_gb":68,"disk_gb":108,"requirement_source":70,"confidence":71},4.74,65.08,{"minimum":110,"production":111,"recommended":112},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":52,"kv_cache_gb":66,"vram_gb":101,"ram_gb":68,"disk_gb":102,"requirement_source":70,"confidence":71},{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":52,"kv_cache_gb":10,"vram_gb":104,"ram_gb":68,"disk_gb":105,"requirement_source":70,"confidence":71},{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":52,"kv_cache_gb":81,"vram_gb":107,"ram_gb":68,"disk_gb":108,"requirement_source":70,"confidence":71},{"minimum":114,"production":116,"recommended":117},{"scenario":63,"context_length":64,"batch_size":42,"weight_gb":52,"kv_cache_gb":66,"vram_gb":115,"ram_gb":68,"disk_gb":102,"requirement_source":70,"confidence":71},2.53,{"scenario":73,"context_length":74,"batch_size":36,"weight_gb":52,"kv_cache_gb":10,"vram_gb":104,"ram_gb":68,"disk_gb":105,"requirement_source":70,"confidence":71},{"scenario":78,"context_length":79,"batch_size":80,"weight_gb":52,"kv_cache_gb":81,"vram_gb":107,"ram_gb":68,"disk_gb":108,"requirement_source":70,"confidence":71},{"recommended":119,"cheapest":137,"performance":146,"min_complexity":163,"cluster_options":178,"no_match":179,"required_vram_gb":107,"candidate_count":180,"usage":181},{"gpu_model":120,"gpu_model_id":121,"gpu_count":42,"total_vram_gb":41,"hour_price":122,"monthly_cost":123,"providers":124,"bandwidth_gbps":126,"fp16_tflops":127,"interconnect_score":128,"parallel_efficiency":42,"score":129,"tier":130,"perf_metric":131,"perf_label":132,"perf_value":133,"workload":16,"cost_metric":134},"RTX 5080",23,0.1339,97.75,[125],"vast",960,113,40,0.5738,"cheapest","tok\u002Fs","19200 tok\u002Fs",19200,{"effective_value":135,"effective_unit":131,"capacity_factor":136},13440,0.7,{"gpu_model":138,"gpu_model_id":139,"gpu_count":42,"total_vram_gb":41,"hour_price":140,"monthly_cost":141,"providers":142,"bandwidth_gbps":43,"fp16_tflops":43,"interconnect_score":128,"parallel_efficiency":42,"score":143,"tier":130,"perf_metric":131,"perf_label":144,"perf_value":43,"workload":16,"cost_metric":145},"Tesla V100",17,0.0289,21.1,[125],0.5686,"0 tok\u002Fs",{"effective_value":43,"effective_unit":131,"capacity_factor":136},{"gpu_model":147,"gpu_model_id":148,"gpu_count":46,"total_vram_gb":149,"hour_price":150,"monthly_cost":151,"providers":152,"bandwidth_gbps":154,"fp16_tflops":155,"interconnect_score":128,"parallel_efficiency":156,"score":157,"tier":158,"perf_metric":131,"perf_label":159,"perf_value":160,"workload":16,"cost_metric":161},"H200 SXM",11,1128,3.5,20440,[153],"lambda",4800,989,0.4,0.2197,"balanced","307200 tok\u002Fs",307200,{"effective_value":162,"effective_unit":131,"capacity_factor":136},215040,{"gpu_model":164,"gpu_model_id":165,"gpu_count":42,"total_vram_gb":166,"hour_price":167,"monthly_cost":168,"providers":169,"bandwidth_gbps":171,"fp16_tflops":172,"interconnect_score":128,"parallel_efficiency":42,"score":173,"tier":130,"perf_metric":131,"perf_label":174,"perf_value":175,"workload":16,"cost_metric":176},"RTX 3090",3,24,0.0622,45.41,[125,170],"tensordock",936,71,0.5484,"18720 tok\u002Fs",18720,{"effective_value":177,"effective_unit":131,"capacity_factor":136},13104,[],false,26,"value",[183,187,191,198,205],{"name":138,"vram_gb":41,"cheapest_hour":140,"cost":184,"provider":125,"estimated_tps":43},{"hourly":140,"daily":185,"monthly":141,"yearly":186},0.69,253.16,{"name":164,"vram_gb":166,"cheapest_hour":167,"cost":188,"provider":125,"estimated_tps":175},{"hourly":167,"daily":189,"monthly":168,"yearly":190},1.49,544.87,{"name":192,"vram_gb":41,"cheapest_hour":193,"cost":194,"provider":125,"estimated_tps":43},"RTX 4070S Ti",0.0678,{"hourly":193,"daily":195,"monthly":196,"yearly":197},1.63,49.49,593.93,{"name":199,"vram_gb":41,"cheapest_hour":200,"cost":201,"provider":125,"estimated_tps":43},"RTX 4080S",0.0685,{"hourly":200,"daily":202,"monthly":203,"yearly":204},1.64,50.01,600.06,{"name":206,"vram_gb":46,"cheapest_hour":207,"cost":208,"provider":125,"estimated_tps":212},"RTX 5060 Ti",0.0719,{"hourly":207,"daily":209,"monthly":210,"yearly":211},1.73,52.49,629.84,8960,{"posts_count":43,"benchmarks_count":43}]