{"id":"W4398457448","doi":"10.7910/dvn/hgihgc/5afaas","title":"study 1 analyses.R","year":2020,"lang":"en","type":"dataset","venue":"Harvard Dataverse","topic":"Reliability and Agreement in Measurement","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Biology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","metaepi_narrow","insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.005209753,0.000471788,0.0009321333,0.0003815509,0.0002336615,0.000805093,0.004285842,0.0002035919,0.08204813],"category_scores_gemma":[0.00843936,0.000335284,0.000348675,0.00110645,0.000117642,0.000404729,0.001578365,0.0006719866,0.7128122],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0001057269,"about_ca_system_score_gemma":0.0002144356,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0003034751,"about_ca_topic_score_gemma":0.0003675477,"domain_scores_codex":[0.9900882,0.001000721,0.001436515,0.00158946,0.005489829,0.0003952384],"domain_scores_gemma":[0.9934803,0.0006935783,0.0006383351,0.00446575,0.0003769077,0.0003450903],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.00004690655,0.0005835979,0.0001505276,0.00001500698,0.0001607151,0.0002522894,0.00007079348,0.00003643989,0.000008934524,0.000002244668,0.9982451,0.0004274154],"study_design_scores_gemma":[0.0005092964,0.0002821933,0.0003573345,0.0000189243,0.0002939031,0.000002434269,0.001485601,0.00001579088,0.000006977154,0.000205864,0.996457,0.000364691],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"dataset","genre_gemma":"dataset","genre_scores_codex":[0.00008424631,8.498526e-7,0.0000663571,0.00004119283,0.001186878,0.0008577271,0.9973587,0.00004041332,0.0003636087],"genre_scores_gemma":[0.0001781554,0.00004902263,0.0001248148,0.0009442011,0.0004057089,0.0000418465,0.9978503,0.00001272428,0.0003932415],"genre_candidate":"dataset","genre_consensus":"dataset","teacher_disagreement_score":0.6307641,"threshold_uncertainty_score":0.999913,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.3457737450820914,"score_gpt":0.4375629867893381,"score_spread":0.09178924170724667,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}