{"id":"W2266911305","doi":"10.1177/0734282915608575","title":"Validation Through Understanding Test-Taking Strategies","year":2015,"lang":"en","type":"article","venue":"Journal of Psychoeducational Assessment","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":9,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of British Columbia","funders":"","keywords":"Test (biology); Psychology; Meaning (existential); Test score; Reading comprehension; Test validity; Task (project management); Congruence (geometry); Cognitive psychology; Reading (process); Social psychology; Psychometrics; Standardized test; Developmental psychology; Mathematics education; Linguistics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.2752618,0.001883589,0.001468947,0.006273388,0.002003058,0.007897194,0.005162615,0.002088665,0.001653804],"category_scores_gemma":[0.5566875,0.0009213156,0.001875432,0.003477038,0.005543749,0.007781155,0.006471965,0.002859771,0.0005831253],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.004084625,"about_ca_system_score_gemma":0.006746037,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004241774,"about_ca_topic_score_gemma":0.003889231,"domain_scores_codex":[0.7056531,0.2320756,0.01879356,0.01154388,0.02877133,0.003162485],"domain_scores_gemma":[0.2924013,0.5454515,0.0285491,0.07077657,0.06088808,0.001933436],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.0006134275,0.0009271622,0.2264991,0.0009981007,0.0005821431,0.0002521336,0.08702599,0.003432963,0.004905226,0.03753417,0.001836186,0.6353934],"study_design_scores_gemma":[0.00079427,0.006707635,0.4588983,0.005124289,0.001152398,0.001448388,0.08248297,0.1279009,0.07184447,0.1615184,0.08111135,0.001016686],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3703654,0.000918685,0.6064883,0.002370681,0.000234093,0.002506784,0.0002265571,0.0006293061,0.0162602],"genre_scores_gemma":[0.7010421,0.0002657862,0.2913867,0.0006181899,0.00009562015,0.004007273,0.0004030922,0.0001738267,0.002007512],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.2752618,"threshold_uncertainty_score":0.8937312,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2924062315474429,"score_gpt":0.4272753511156951,"score_spread":0.1348691195682522,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}