{"id":"W2146073079","doi":"","title":"TIMSS Matematik Test Sonuçlarının Amerika ve Türkiye Arasında Karşılaştırılabilirliği","year":2014,"lang":"tr","type":"article","venue":"EĞİTİM VE BİLİM","topic":"Educational Assessment and Pedagogy","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of British Columbia; Ottawa Hospital","funders":"","keywords":"Comparability; Test (biology); Mathematics; Exploratory factor analysis; Statistics; Scale (ratio); Differential item functioning; Mathematics education; Psychology; Item response theory; Geography; Combinatorics; Psychometrics; Cartography","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaepi_narrow","sts","insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.002769012,0.0009225875,0.001160729,0.0003522368,0.001574817,0.0008419073,0.001632293,0.0006807093,0.008733577],"category_scores_gemma":[0.001175615,0.0009223172,0.0005052286,0.001450984,0.001182704,0.0008727408,0.000324854,0.0008788597,0.01011748],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005552375,"about_ca_system_score_gemma":0.001441325,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00299412,"about_ca_topic_score_gemma":0.001957396,"domain_scores_codex":[0.9927086,0.0008853928,0.001346147,0.001442206,0.001823568,0.001794055],"domain_scores_gemma":[0.9939952,0.00256856,0.0009250854,0.001038038,0.0006551414,0.0008179099],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0002151794,0.003998688,0.2984936,0.001443293,0.0008102486,0.00005259462,0.09229672,0.0000836542,0.00392151,0.1540818,0.4001115,0.04449124],"study_design_scores_gemma":[0.001272498,0.0008869682,0.06867719,0.0004801181,0.0003573741,0.00001740844,0.02587887,0.001072493,0.0004947937,0.01442982,0.8843967,0.002035711],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6200379,0.00219738,0.0003334708,0.03146969,0.01015251,0.002119868,0.0003628418,0.0005841174,0.3327422],"genre_scores_gemma":[0.9407938,0.000690492,0.001283459,0.002075536,0.005989311,0.0001831722,0.0001995547,0.0001369083,0.04864779],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.4842853,"threshold_uncertainty_score":0.999725,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03549380525329934,"score_gpt":0.3716038008047365,"score_spread":0.3361099955514372,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}