{"id":"W2148427497","doi":"10.7202/1025007ar","title":"The versatility of generalizability theory as a tool for exploring and controlling measurement error","year":2014,"lang":"en","type":"article","venue":"Mesure et évaluation en éducation","topic":"Educational Assessment and Improvement","field":"Decision Sciences","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Generalizability theory; Computer science; Sample (material); Observational error; Estimation; Item response theory; Econometrics; Statistics; Psychology; Mathematics; Psychometrics; Engineering","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.3848927,0.003383557,0.005079138,0.01616807,0.003035214,0.009175404,0.005206907,0.004886488,0.003653033],"category_scores_gemma":[0.6076586,0.001842351,0.005214983,0.01338895,0.02290763,0.01590542,0.01293615,0.01079223,0.0005838589],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.005327112,"about_ca_system_score_gemma":0.00651439,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003994287,"about_ca_topic_score_gemma":0.002619947,"domain_scores_codex":[0.4676945,0.4684818,0.01327714,0.01588184,0.03327956,0.001385136],"domain_scores_gemma":[0.1882224,0.7281266,0.01648275,0.05663581,0.009755543,0.0007770486],"domain_codex":"methods","domain_gemma":"methods","domain_candidate":"methods","domain_consensus":"methods","study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001785987,0.0002153558,0.02199212,0.003117086,0.00364156,0.0003625227,0.009121269,0.01348394,0.0007522693,0.7860259,0.002302291,0.158807],"study_design_scores_gemma":[0.0001302385,0.0005587695,0.008957531,0.00145114,0.0007620198,0.0003730349,0.001387471,0.02349236,0.001126916,0.952397,0.00919844,0.0001649621],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.005829428,0.001810336,0.9774935,0.003674991,0.0002982026,0.001149798,0.0001872563,0.0003301821,0.009226249],"genre_scores_gemma":[0.4031219,0.002400039,0.5826563,0.001729837,0.0009234191,0.007530778,0.0002332499,0.0002993048,0.001105215],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.6151073,"threshold_uncertainty_score":0.7585368,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3700201145714859,"score_gpt":0.4653012208777185,"score_spread":0.0952811063062326,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}