{"id":"W4200090177","doi":"10.21203/rs.3.rs-1117982/v1","title":"Probing the Effect of Selection Bias on Generalization: A Thought Experiment","year":2021,"lang":"en","type":"preprint","venue":"Research Square","topic":"Bayesian Modeling and Causal Inference","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"Huawei Technologies (Canada); York University","funders":"","keywords":"Generalization; Set (abstract data type); Computer science; Focus (optics); Artificial intelligence; Task (project management); Range (aeronautics); Domain (mathematical analysis); Machine learning; Point (geometry); Population; Cognitive psychology; Psychology; Epistemology; Mathematics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0456153,0.0008993791,0.0009253385,0.0005296011,0.001222163,0.00256091,0.002484439,0.002713341,0.01165629],"category_scores_gemma":[0.1654872,0.0005627593,0.001445428,0.0007010514,0.006911692,0.006292856,0.003057959,0.005570916,0.001560413],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001222355,"about_ca_system_score_gemma":0.0008073063,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0009063731,"about_ca_topic_score_gemma":0.0004519493,"domain_scores_codex":[0.9789722,0.01492028,0.0008865881,0.00267189,0.002113549,0.0004354281],"domain_scores_gemma":[0.7157325,0.2406609,0.01085104,0.02498722,0.005953429,0.00181479],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0331306,0.014419,0.05242318,0.004050259,0.002496186,0.001809112,0.03327513,0.02059726,0.06193016,0.4537474,0.05881636,0.2633054],"study_design_scores_gemma":[0.01023341,0.01606231,0.02853257,0.0006879764,0.001103371,0.001242326,0.004135862,0.06552093,0.03869747,0.7749367,0.05835027,0.0004968047],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"empirical","genre_gemma":"methods","genre_scores_codex":[0.7579576,0.001728039,0.1434144,0.02866041,0.002024091,0.002988228,0.001466783,0.0008390116,0.06092154],"genre_scores_gemma":[0.8988653,0.0004441893,0.0815884,0.00987157,0.0009346254,0.00230481,0.000489563,0.0003389004,0.005162644],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.0456153,"threshold_uncertainty_score":0.2412397,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.1153445373749779,"score_gpt":0.4022892274625641,"score_spread":0.2869446900875862,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}