{"id":"W2136085695","doi":"10.5539/ies.v6n6p185","title":"Rasch Model Analysis on the Effectiveness of Early Evaluation Questions as a Benchmark for New Students Ability","year":2013,"lang":"en","type":"article","venue":"International Education Studies","topic":"Psychometric Methodologies and Testing","field":"Decision Sciences","cited_by":5,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Rasch model; Cronbach's alpha; Psychology; Benchmark (surveying); Mathematics education; Relevance (law); Item analysis; Measure (data warehouse); Polytomous Rasch model; Item response theory; Psychometrics; Computer science; Developmental psychology; Data mining","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.05604182,0.001341809,0.001289986,0.003420197,0.0008438452,0.002127327,0.0008167392,0.0009140056,0.001852899],"category_scores_gemma":[0.1651961,0.0003266555,0.002376026,0.002386106,0.0009413164,0.002099935,0.001157585,0.001460037,0.0005371227],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008643693,"about_ca_system_score_gemma":0.001397142,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004709975,"about_ca_topic_score_gemma":0.003182042,"domain_scores_codex":[0.9473546,0.04004133,0.002135997,0.002386277,0.007402969,0.0006787915],"domain_scores_gemma":[0.7172711,0.2517883,0.005930319,0.007635755,0.01671266,0.0006618876],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"observational","study_design_gemma":"observational","study_design_scores_codex":[0.003518058,0.001777,0.5557774,0.0009578754,0.002958648,0.0003046354,0.009065775,0.07043094,0.005312294,0.00678654,0.003223618,0.3398872],"study_design_scores_gemma":[0.0002123709,0.009981045,0.4729142,0.0005777052,0.001138037,0.0003990401,0.005294594,0.4853745,0.01217802,0.006680113,0.004914805,0.000335742],"study_design_candidate":"observational","study_design_consensus":"observational","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8738406,0.0007505858,0.1143798,0.000349015,0.0001679574,0.0009696364,0.0004981299,0.00064475,0.00839966],"genre_scores_gemma":[0.9743973,0.0001113189,0.02407124,0.00003051471,0.00002067428,0.0003671646,0.0002824075,0.00004104993,0.0006783524],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.05604182,"threshold_uncertainty_score":0.296381,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.5771839040710531,"score_gpt":0.6243995894698064,"score_spread":0.04721568539875332,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}