{"id":"W1572437522","doi":"","title":"A Needs-Based Approach to the Evaluation of the Spoken Language Ability of International Teaching Assistants","year":2002,"lang":"en","type":"article","venue":"DOAJ (DOAJ: Directory of Open Access Journals)","topic":"EFL/ESL Teaching and Learning","field":"Arts and Humanities","cited_by":10,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Test (biology); Language assessment; Language proficiency; Graduate students; Reliability (semiconductor); Psychology; Task (project management); Language education; Test validity; Computer science; Mathematics education; Pedagogy; Psychometrics; Engineering","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01410992,0.0003955688,0.0003928729,0.005286216,0.001526939,0.00192576,0.00103259,0.0007726332,0.002625061],"category_scores_gemma":[0.03751486,0.0002742556,0.000422672,0.00174758,0.001614439,0.001768544,0.002756793,0.001003429,0.0003584283],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.003368951,"about_ca_system_score_gemma":0.003797853,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007250144,"about_ca_topic_score_gemma":0.009466225,"domain_scores_codex":[0.9840528,0.01028547,0.001023757,0.0003008362,0.003641539,0.0006956542],"domain_scores_gemma":[0.9750366,0.01129626,0.001747982,0.0008379495,0.009355704,0.001725459],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.001175397,0.004583669,0.4041124,0.001032547,0.00009064373,0.001225423,0.2304039,0.003807386,0.009736654,0.01565983,0.002972483,0.3251997],"study_design_scores_gemma":[0.0002291511,0.004747741,0.4431762,0.0004039565,0.0001161467,0.00102738,0.4814494,0.02497878,0.01371728,0.01729428,0.0126552,0.0002045947],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9804657,0.00007761639,0.006256291,0.00039132,0.00001716272,0.0007793546,0.00008195252,0.00003701196,0.01189364],"genre_scores_gemma":[0.9905205,0.00004450519,0.007499521,0.00004184953,0.000007783367,0.0006860191,0.00004366054,0.000005157223,0.001151025],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.01410992,"threshold_uncertainty_score":0.07462132,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.4162669563386445,"score_gpt":0.538539862983312,"score_spread":0.1222729066446676,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}