{"id":"W4399100676","doi":"10.23941/ejpe.v17i1.849","title":"The Challenge of Choosing Well","year":2024,"lang":"en","type":"article","venue":"Erasmus Journal for Philosophy and Economics","topic":"Evaluation and Performance Assessment","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"","keywords":"Computer science","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003606907,0.00007208339,0.0001389369,0.0001099889,0.0004373207,0.0005229498,0.0002291176,0.00003451835,0.0001820923],"category_scores_gemma":[0.0001104628,0.00004180718,0.000123338,0.00007053985,0.00006343571,0.0004004958,0.00003132793,0.000146725,0.00004073408],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00002759576,"about_ca_system_score_gemma":0.00009700602,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":6.090339e-7,"about_ca_topic_score_gemma":0.000006052793,"domain_scores_codex":[0.9990031,0.00003519033,0.0005339435,0.0001458026,0.0001587012,0.0001233378],"domain_scores_gemma":[0.9986589,0.0008523109,0.0001672532,0.0001486871,0.0001046908,0.00006816102],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0001047657,0.00002823854,0.0003388867,0.00003063597,0.00009645034,0.000001416072,0.001518928,0.001250264,0.00001846848,0.6895122,0.005967312,0.3011324],"study_design_scores_gemma":[0.000171774,0.00009726074,0.0001038637,0.00002099651,0.000008015487,0.0000212761,0.0001615805,0.02419805,0.00003777455,0.6570899,0.3180414,0.00004799477],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.5865501,0.02292906,0.003272922,0.1706848,0.0151557,0.0006523234,0.0000829839,0.00003620105,0.2006359],"genre_scores_gemma":[0.9933264,0.004843723,0.0002385093,0.0001662774,0.0007609388,0.000004119541,8.916013e-7,0.000007249404,0.0006519408],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.4067762,"threshold_uncertainty_score":0.5042817,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.2090684220383123,"score_gpt":0.4478009860038947,"score_spread":0.2387325639655823,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}