{"id":"W3196166069","doi":"10.33774/chemrxiv-2021-7m4tw","title":"Response process validity evidence in chemistry education research","year":2021,"lang":"en","type":"preprint","venue":"","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Ottawa","funders":"","keywords":"Process (computing); Cognition; Psychology; Management science; Computer science; Applied psychology; Data science; Engineering; Neuroscience","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.6945845,0.001817613,0.002887075,0.01070715,0.005782409,0.01617708,0.006308811,0.009538685,0.01825407],"category_scores_gemma":[0.9064366,0.002511223,0.005978423,0.01344162,0.02264572,0.01851406,0.01205838,0.009864502,0.003591911],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.01149848,"about_ca_system_score_gemma":0.02201089,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004073864,"about_ca_topic_score_gemma":0.003354022,"domain_scores_codex":[0.1803973,0.5656902,0.08263551,0.0284606,0.1382635,0.004552877],"domain_scores_gemma":[0.02695165,0.8483456,0.02824191,0.04634461,0.04921635,0.0008998433],"domain_codex":"methods","domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.00211768,0.0009185119,0.07469164,0.05292312,0.003974988,0.000473098,0.04397523,0.002932188,0.001823928,0.3762336,0.01833187,0.4216042],"study_design_scores_gemma":[0.001488047,0.002732647,0.07919694,0.125629,0.003534099,0.001305983,0.02073615,0.01083643,0.01498939,0.4841444,0.2547613,0.0006456194],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.08395591,0.06706788,0.5544153,0.0878014,0.007510557,0.01343957,0.003070226,0.0006524392,0.1820867],"genre_scores_gemma":[0.6882344,0.01806029,0.2379348,0.02229946,0.002204555,0.022961,0.002967329,0.001303782,0.004034454],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.6945845,"threshold_uncertainty_score":0.3766317,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.7982323329816771,"score_gpt":0.6964298082778361,"score_spread":0.101802524703841,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}