{"id":"W2123012719","doi":"10.1177/026553220101800302","title":"Examining dialogue: another approach to content specification and to validating inferences drawn from test scores","year":2001,"lang":"en","type":"article","venue":"Language Testing","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":162,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Construct (python library); Psychology; Test (biology); Sociocultural evolution; Point (geometry); Cognition; Cognitive psychology; Content (measure theory); Mathematics education; Inference; Construct validity; Linguistics; Natural language processing; Social psychology; Computer science; Artificial intelligence; Psychometrics; Developmental psychology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.161656,0.002729091,0.002763749,0.02730363,0.003611698,0.01562342,0.005862653,0.004393897,0.004174567],"category_scores_gemma":[0.4097496,0.000802374,0.002223392,0.01508754,0.009089683,0.01691436,0.009106664,0.005165908,0.001536712],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003753197,"about_ca_system_score_gemma":0.006489226,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00533489,"about_ca_topic_score_gemma":0.003941027,"domain_scores_codex":[0.7846398,0.1580065,0.01798476,0.0113404,0.02527799,0.002750428],"domain_scores_gemma":[0.3497685,0.4996819,0.03248017,0.05257063,0.06253923,0.002959656],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.00109646,0.001500914,0.1126694,0.002429954,0.0007175228,0.001087532,0.109102,0.005130205,0.01836386,0.1303187,0.004889062,0.6126944],"study_design_scores_gemma":[0.0005741596,0.005203129,0.1279792,0.003305797,0.001004626,0.002155275,0.110943,0.09504624,0.0782926,0.459444,0.1147174,0.001334486],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1208789,0.0003642289,0.8447046,0.002643687,0.0002812423,0.00281024,0.001271952,0.001939338,0.02510589],"genre_scores_gemma":[0.3306119,0.0001701277,0.6593288,0.0008451003,0.0001841323,0.004014475,0.001191895,0.0004317374,0.003221876],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.161656,"threshold_uncertainty_score":0.8549289,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2440808435611483,"score_gpt":0.2751833738551075,"score_spread":0.03110253029395921,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}