{"id":"W2110105050","doi":"10.5430/wje.v2n4p102","title":"Dynamic Assessment of EFL Reading: Revealing Hidden Aspects at Different Proficiency Levels","year":2012,"lang":"en","type":"article","venue":"World Journal of Education","topic":"Educational and Psychological Assessments","field":"Psychology","cited_by":17,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Reading (process); Psychology; Mathematics education; Test (biology); Language proficiency; Developmental psychology; Cognitive psychology; Linguistics","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0006633785,0.0001528243,0.0002867299,0.0003348037,0.00008902237,0.00001536354,0.0002561572,0.00006128209,0.002809522],"category_scores_gemma":[0.00004963656,0.0001150644,0.0001278117,0.0003385602,0.00004969919,0.0001786592,0.00003118042,0.0002833474,0.00003733411],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004238132,"about_ca_system_score_gemma":0.0001679266,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00001028575,"about_ca_topic_score_gemma":0.000005576156,"domain_scores_codex":[0.9982137,0.0002157943,0.0007317663,0.0001612424,0.0003797388,0.0002977703],"domain_scores_gemma":[0.9981127,0.0002210307,0.0009859481,0.0002347734,0.0002484531,0.0001970915],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.00007559069,0.006483221,0.8667179,0.00007283114,0.0001922808,0.000001875262,0.00415956,0.000004207538,0.01510363,0.06746065,0.01392272,0.02580548],"study_design_scores_gemma":[0.0002793446,0.0001553514,0.9918072,0.0001362228,0.00006008908,0.00007178553,0.001032995,0.000003182073,0.0002681973,0.003444731,0.002619768,0.0001211939],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.9347313,0.001023207,0.0002800682,0.001892795,0.006544033,0.0001833739,0.00000495704,0.000007356927,0.05533293],"genre_scores_gemma":[0.9824055,0.00001871906,0.002867793,0.0002270458,0.0005493887,0.00001978827,0.00001269092,0.00001461317,0.01388447],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.1250892,"threshold_uncertainty_score":0.9981021,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0642124569722308,"score_gpt":0.438874863380181,"score_spread":0.3746624064079502,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}