{"id":"W4200191876","doi":"10.1109/embc46164.2021.9629697","title":"Machine Learning Model Validation for Early Stage Studies with Small Sample Sizes","year":2021,"lang":"en","type":"article","venue":"2021 43rd Annual International Conference of the IEEE Engineering in Medicine &amp; Biology Society (EMBC)","topic":"Advanced Causal Inference Techniques","field":"Mathematics","cited_by":39,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of New Brunswick","funders":"","keywords":"Overfitting; Feature selection; Feature (linguistics); Cross-validation; Selection (genetic algorithm); Sample size determination; Model selection; Sample (material)","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.2605956,0.002150703,0.002066172,0.001954344,0.00214398,0.00305663,0.00345597,0.002690391,0.003724406],"category_scores_gemma":[0.4908872,0.0009975438,0.003354784,0.001599778,0.003007654,0.003936213,0.003298916,0.005471942,0.0008118679],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001719129,"about_ca_system_score_gemma":0.005408662,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002525151,"about_ca_topic_score_gemma":0.003366166,"domain_scores_codex":[0.8673836,0.1156173,0.005766044,0.004877102,0.005490765,0.0008652745],"domain_scores_gemma":[0.4899657,0.4344595,0.01218193,0.04372811,0.01874103,0.0009236715],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.006190356,0.002250513,0.1144104,0.006841378,0.006044276,0.001395275,0.00443129,0.2243793,0.02182148,0.1105928,0.01663121,0.4850117],"study_design_scores_gemma":[0.001203236,0.006570597,0.04447106,0.003546879,0.001471571,0.0007619713,0.001124608,0.6648002,0.03414985,0.1935194,0.04795777,0.0004228212],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01837431,0.0007783807,0.9761486,0.000580192,0.000195574,0.002030658,0.0003440302,0.0004740122,0.001074311],"genre_scores_gemma":[0.2826722,0.0004945141,0.7027878,0.001035084,0.0001020064,0.01063713,0.001059375,0.0003853485,0.0008265718],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.7394043,"threshold_uncertainty_score":0.9118172,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2296422583971432,"score_gpt":0.4107561751219546,"score_spread":0.1811139167248114,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}