[{"data":1,"prerenderedAt":-1},["ShallowReactive",2],{"model-qwen3-asr-1-7b":3},{"id":4,"slug":5,"name":6,"hf_id":6,"organization":7,"architecture":8,"architecture_config":9,"parameters":10,"parameters_label":11,"context_length":12,"license":12,"downloads":13,"likes":14,"status":15,"description":12,"huggingface_url":16,"created_at":17,"seo":18,"category":21,"family":24,"score":26,"variants":27,"quantizations":28,"requirements":55,"gpu_recommendations":125,"fitting_gpus":191,"community":225},512,"qwen3-asr-1-7b","Qwen\u002FQwen3-ASR-1.7B","Qwen","qwen3_asr",[],2349217408,"2.3B",null,4206898,1066,"published","https:\u002F\u002Fhuggingface.co\u002FQwen\u002FQwen3-ASR-1.7B","2026-08-15T12:14:53+00:00",{"title":6,"description":19,"image":12,"robots":20,"canonical_url":12},"AI model deployment & compute intelligence — compare AI models, GPU requirements, benchmarks, cloud GPU rental pricing and deployment cost analysis for DeepSeek, Qwen, Llama, Wan and more.","index, follow",{"name":22,"slug":23},"Audio","audio",{"name":7,"slug":25},"qwen",0.0594,[],[29,34,39,44,48,52],{"format":30,"bits":31,"size_factor":32,"quality_loss":33},"AWQ",4,0.263,0.03,{"format":35,"bits":36,"size_factor":37,"quality_loss":38},"FP16",16,1,0,{"format":40,"bits":41,"size_factor":42,"quality_loss":43},"FP8",8,0.5,0.01,{"format":45,"bits":31,"size_factor":46,"quality_loss":47},"GGUF",0.25,0.05,{"format":49,"bits":31,"size_factor":50,"quality_loss":51},"GPTQ",0.275,0.04,{"format":53,"bits":41,"size_factor":42,"quality_loss":54},"INT8",0.02,{"FP16":56,"FP8":77,"INT8":88,"AWQ":92,"GPTQ":103,"GGUF":114},{"minimum":57,"production":66,"recommended":71},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":60,"kv_cache_gb":38,"vram_gb":61,"ram_gb":62,"disk_gb":63,"requirement_source":64,"confidence":65},"minimum",4096,4.38,13.34,32,45.83,"inferred",0.95,{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":60,"kv_cache_gb":38,"vram_gb":69,"ram_gb":62,"disk_gb":70,"requirement_source":64,"confidence":65},"production",16384,16.35,110.83,{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":60,"kv_cache_gb":38,"vram_gb":75,"ram_gb":62,"disk_gb":76,"requirement_source":64,"confidence":65},"recommended",8192,2,14.52,71.83,{"minimum":78,"production":82,"recommended":85},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":79,"kv_cache_gb":38,"vram_gb":80,"ram_gb":62,"disk_gb":81,"requirement_source":64,"confidence":65},2.19,8.04,42.41,{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":79,"kv_cache_gb":38,"vram_gb":83,"ram_gb":62,"disk_gb":84,"requirement_source":64,"confidence":65},10.58,107.41,{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":79,"kv_cache_gb":38,"vram_gb":86,"ram_gb":62,"disk_gb":87,"requirement_source":64,"confidence":65},8.99,68.41,{"minimum":89,"production":90,"recommended":91},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":79,"kv_cache_gb":38,"vram_gb":80,"ram_gb":62,"disk_gb":81,"requirement_source":64,"confidence":65},{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":79,"kv_cache_gb":38,"vram_gb":83,"ram_gb":62,"disk_gb":84,"requirement_source":64,"confidence":65},{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":79,"kv_cache_gb":38,"vram_gb":86,"ram_gb":62,"disk_gb":87,"requirement_source":64,"confidence":65},{"minimum":93,"production":97,"recommended":100},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":94,"kv_cache_gb":38,"vram_gb":95,"ram_gb":62,"disk_gb":96,"requirement_source":64,"confidence":65},1.13,5.5,40.77,{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":94,"kv_cache_gb":38,"vram_gb":98,"ram_gb":62,"disk_gb":99,"requirement_source":64,"confidence":65},7.8,105.77,{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":94,"kv_cache_gb":38,"vram_gb":101,"ram_gb":62,"disk_gb":102,"requirement_source":64,"confidence":65},6.32,66.77,{"minimum":104,"production":108,"recommended":111},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":105,"kv_cache_gb":38,"vram_gb":106,"ram_gb":62,"disk_gb":107,"requirement_source":64,"confidence":65},1.15,5.53,40.79,{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":105,"kv_cache_gb":38,"vram_gb":109,"ram_gb":62,"disk_gb":110,"requirement_source":64,"confidence":65},7.83,105.79,{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":105,"kv_cache_gb":38,"vram_gb":112,"ram_gb":62,"disk_gb":113,"requirement_source":64,"confidence":65},6.36,66.79,{"minimum":115,"production":119,"recommended":122},{"scenario":58,"context_length":59,"batch_size":37,"weight_gb":116,"kv_cache_gb":38,"vram_gb":117,"ram_gb":62,"disk_gb":118,"requirement_source":64,"confidence":65},1.18,5.6,40.83,{"scenario":67,"context_length":68,"batch_size":31,"weight_gb":116,"kv_cache_gb":38,"vram_gb":120,"ram_gb":62,"disk_gb":121,"requirement_source":64,"confidence":65},7.9,105.83,{"scenario":72,"context_length":73,"batch_size":74,"weight_gb":116,"kv_cache_gb":38,"vram_gb":123,"ram_gb":62,"disk_gb":124,"requirement_source":64,"confidence":65},6.43,66.83,{"recommended":126,"cheapest":145,"performance":158,"min_complexity":173,"cluster_options":187,"no_match":188,"required_vram_gb":101,"candidate_count":189,"usage":190},{"gpu_model":127,"gpu_model_id":128,"gpu_count":37,"total_vram_gb":129,"hour_price":130,"monthly_cost":131,"providers":132,"bandwidth_gbps":134,"fp16_tflops":135,"interconnect_score":136,"parallel_efficiency":37,"score":137,"tier":138,"perf_metric":139,"perf_label":140,"perf_value":141,"workload":12,"cost_metric":142},"RTX 3080 Ti",36,12,0.1089,79.5,[133],"vast",912,68,40,0.6353,"cheapest","RTF","RTF 0.39",0.39,{"effective_value":143,"effective_unit":139,"capacity_factor":144},0.3,0.7,{"gpu_model":146,"gpu_model_id":147,"gpu_count":37,"total_vram_gb":148,"hour_price":149,"monthly_cost":150,"providers":151,"bandwidth_gbps":152,"fp16_tflops":153,"interconnect_score":136,"parallel_efficiency":37,"score":154,"tier":138,"perf_metric":139,"perf_label":155,"perf_value":156,"workload":12,"cost_metric":157},"RTX 3080",35,10,0.0806,58.84,[133],760,60,0.6627,"RTF 0.46",0.46,{"effective_value":143,"effective_unit":139,"capacity_factor":144},{"gpu_model":159,"gpu_model_id":74,"gpu_count":41,"total_vram_gb":160,"hour_price":161,"monthly_cost":162,"providers":163,"bandwidth_gbps":165,"fp16_tflops":166,"interconnect_score":136,"parallel_efficiency":167,"score":168,"tier":169,"perf_metric":139,"perf_label":170,"perf_value":171,"workload":12,"cost_metric":172},"RTX 5090",256,0.4424,2583.62,[133,164],"runpod",1792,209,0.4,0.3471,"balanced","RTF 0.06",0.06,{"effective_value":38,"effective_unit":139,"capacity_factor":144},{"gpu_model":174,"gpu_model_id":175,"gpu_count":37,"total_vram_gb":176,"hour_price":177,"monthly_cost":178,"providers":179,"bandwidth_gbps":181,"fp16_tflops":182,"interconnect_score":136,"parallel_efficiency":37,"score":183,"tier":138,"perf_metric":139,"perf_label":184,"perf_value":185,"workload":12,"cost_metric":186},"RTX 3090",3,24,0.1341,97.89,[133,180],"tensordock",936,71,0.5649,"RTF 0.38",0.38,{"effective_value":143,"effective_unit":139,"capacity_factor":144},[],false,17,"value",[192,197,204,212,217],{"name":146,"vram_gb":148,"cheapest_hour":149,"cost":193,"provider":133,"estimated_tps":196},{"hourly":149,"daily":194,"monthly":150,"yearly":195},1.93,706.06,672.6,{"name":198,"vram_gb":36,"cheapest_hour":199,"cost":200,"provider":133,"estimated_tps":38},"Tesla V100",0.0887,{"hourly":199,"daily":201,"monthly":202,"yearly":203},2.13,64.75,777.01,{"name":205,"vram_gb":41,"cheapest_hour":206,"cost":207,"provider":133,"estimated_tps":211},"RTX 4060 Ti",0.1081,{"hourly":206,"daily":208,"monthly":209,"yearly":210},2.59,78.91,946.96,254.9,{"name":127,"vram_gb":129,"cheapest_hour":130,"cost":213,"provider":133,"estimated_tps":216},{"hourly":130,"daily":214,"monthly":131,"yearly":215},2.61,953.96,807.1,{"name":218,"vram_gb":41,"cheapest_hour":219,"cost":220,"provider":133,"estimated_tps":224},"RTX 5060 Ti",0.1222,{"hourly":219,"daily":221,"monthly":222,"yearly":223},2.93,89.21,1070.47,396.5,{"posts_count":38,"benchmarks_count":38}]