{"id":"W4416620164","doi":"10.48550/arxiv.2510.20519","title":"Metis-HOME: Hybrid Optimized Mixture-of-Experts for Multimodal Reasoning","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Constraint Satisfaction and Optimization","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Structuring; Inference; Focus (optics); Field (mathematics); Simple (philosophy); Key (lock); Reversing; Reasoning system","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow"],"consensus_categories":[],"category_scores_codex":[0.001015832,0.0009643091,0.001510457,0.0006638658,0.0005522888,0.0003039671,0.001774929,0.0005905442,0.0004366529],"category_scores_gemma":[0.0009475442,0.001063371,0.001004632,0.0007168886,0.000331013,0.0006210259,0.001588816,0.000786971,0.00002922557],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0002920555,"about_ca_system_score_gemma":0.001205913,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00021059,"about_ca_topic_score_gemma":0.00002304605,"domain_scores_codex":[0.9941378,0.000382818,0.001783993,0.002174746,0.0006235577,0.0008970859],"domain_scores_gemma":[0.9943222,0.0009511377,0.001310325,0.001935944,0.001135227,0.0003451938],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001379509,0.001213025,0.1243256,0.002135353,0.002169532,0.00007443153,0.007106731,0.3977288,0.002402538,0.01034754,0.004514351,0.4466026],"study_design_scores_gemma":[0.004652574,0.0001793291,0.02264534,0.001434452,0.000308078,0.00003908723,0.0001680843,0.9511243,0.0122117,0.0004656198,0.00529713,0.001474306],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.06293106,0.000736701,0.9240094,0.001357573,0.006723369,0.002314736,0.0002032096,0.0002930523,0.001430917],"genre_scores_gemma":[0.5594334,0.0008288612,0.4361617,0.0004359982,0.0002863256,0.0003751719,0.0001745289,0.00005003505,0.002254022],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.5533955,"threshold_uncertainty_score":0.9991816,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02726740771065094,"score_gpt":0.2866986835762741,"score_spread":0.2594312758656232,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}