{"id":"W4416159334","doi":"10.48550/arxiv.2511.07417","title":"Language Generation with Infinite Contamination","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"Machine Learning and Algorithms","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada","keywords":"Robustness (evolution); Countable set; Generator (circuit theory); Limit (mathematics); Oracle; Fraction (chemistry); Minimax; Third generation","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005755926,0.001348034,0.00157859,0.001190077,0.001469111,0.002890223,0.003156072,0.002927881,0.004604244],"category_scores_gemma":[0.06119603,0.0008916202,0.001569416,0.001045322,0.005748401,0.008369111,0.006563935,0.005690595,0.001284385],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002541692,"about_ca_system_score_gemma":0.001304874,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001072545,"about_ca_topic_score_gemma":0.000585114,"domain_scores_codex":[0.9908388,0.003491464,0.0003556562,0.002209062,0.002115893,0.0009890076],"domain_scores_gemma":[0.9301679,0.05030258,0.004634331,0.01086439,0.002198977,0.001831887],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0007767831,0.0003889923,0.008472426,0.0004250007,0.0002028942,0.001171016,0.001721031,0.4333163,0.009925952,0.4993929,0.003536384,0.04067032],"study_design_scores_gemma":[0.00005974762,0.0001733808,0.0005116417,0.00004232468,0.00003110237,0.0003279982,0.0001126598,0.5914443,0.004356696,0.4009004,0.001996546,0.00004318165],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1687526,0.0009263991,0.8110125,0.003267741,0.00008335103,0.0002024519,0.000634629,0.001411234,0.01370905],"genre_scores_gemma":[0.9400578,0.0003143196,0.05123331,0.0007883409,0.0001391556,0.0003119714,0.0005678542,0.0002780471,0.006309043],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.005755926,"threshold_uncertainty_score":0.03044063,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.02631623567508679,"score_gpt":0.2794779266072273,"score_spread":0.2531616909321405,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}