{"id":"W1985695119","doi":"10.1111/j.1745-3992.2004.tb00164.x","title":"Avoiding Misconception, Misuse, and Missed Opportunities: The Collection of Verbal Reports in Educational Achievement Testing","year":2004,"lang":"en","type":"article","venue":"Educational Measurement Issues and Practice","topic":"Educational and Psychological Assessments","field":"Psychology","cited_by":121,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Psychology; Cognition; Data collection; Trustworthiness; Nonverbal communication; Test (biology); Cognitive psychology; Developmental psychology; Applied psychology; Social psychology; Social science","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":["metaresearch"],"category_scores_codex":[0.509223,0.001489825,0.00176673,0.009384218,0.006096591,0.01644592,0.007318873,0.009632415,0.0005064935],"category_scores_gemma":[0.7473314,0.002464328,0.0010475,0.007213284,0.0706352,0.02003948,0.01101502,0.0172739,0.0009481451],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.006880463,"about_ca_system_score_gemma":0.01233807,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005293658,"about_ca_topic_score_gemma":0.00732099,"domain_scores_codex":[0.2746556,0.602209,0.04299914,0.006105453,0.07232972,0.001701086],"domain_scores_gemma":[0.1067892,0.7423745,0.05484158,0.04332184,0.05040879,0.002264001],"domain_codex":"methods","domain_gemma":"methods","domain_candidate":"methods","domain_consensus":"methods","study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0004078995,0.0002019118,0.03218332,0.005135125,0.0003016947,0.002015515,0.3569099,0.000807889,0.002043923,0.1000795,0.04669856,0.4532148],"study_design_scores_gemma":[0.0002324158,0.001456548,0.03814649,0.05722392,0.000724687,0.01426067,0.235245,0.009047845,0.01540563,0.308111,0.3188248,0.00132098],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"commentary","genre_gemma":"methods","genre_scores_codex":[0.1021295,0.07097492,0.294193,0.5039759,0.01142938,0.001014647,0.0002338584,0.001072325,0.0149766],"genre_scores_gemma":[0.6075315,0.0269433,0.201191,0.1461435,0.01102038,0.002500386,0.0001627787,0.0007973716,0.003709642],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.490777,"threshold_uncertainty_score":0.6052154,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3226938867206189,"score_gpt":0.4332718452277097,"score_spread":0.1105779585070908,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}