{"id":"W2888454927","doi":"10.1037/met0000176","title":"A cautionary note on the finite sample behavior of maximal reliability.","year":2018,"lang":"en","type":"article","venue":"Psychological Methods","topic":"Reliability and Maintenance Optimization","field":"Engineering","cited_by":26,"is_retracted":false,"has_abstract":true,"ca_institutions":"Public Safety Canada","funders":"Aalto-Yliopisto; Academy of Finland","keywords":"Reliability (semiconductor); Sample (material); Statistics; Mathematics; Psychology; Econometrics; Applied mathematics; Physics; Thermodynamics","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.1103604,0.001674603,0.002343252,0.002421516,0.00322968,0.005442529,0.008788242,0.005244973,0.004645648],"category_scores_gemma":[0.4194961,0.0008000577,0.002185091,0.003883711,0.01558027,0.007876874,0.003740069,0.03080079,0.003149368],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002660002,"about_ca_system_score_gemma":0.003574376,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01013897,"about_ca_topic_score_gemma":0.01351829,"domain_scores_codex":[0.9012196,0.05795336,0.008918181,0.01288619,0.01811157,0.0009110235],"domain_scores_gemma":[0.5504087,0.3915319,0.008944341,0.02350299,0.02412763,0.00148449],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004725113,0.0001161478,0.01199521,0.001683599,0.0005324279,0.001648366,0.0106315,0.002354444,0.001531847,0.371765,0.4815527,0.1157161],"study_design_scores_gemma":[0.000231975,0.0003297259,0.01576999,0.003202826,0.0003108364,0.002775565,0.002921083,0.01336097,0.004752382,0.5926485,0.3631836,0.0005124563],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"commentary","genre_gemma":"empirical","genre_scores_codex":[0.01479848,0.03109908,0.2870837,0.5910544,0.0387486,0.0004781273,0.002088042,0.002204279,0.03244539],"genre_scores_gemma":[0.2459541,0.009911564,0.3447807,0.3418185,0.02827161,0.002099856,0.0004893505,0.001356277,0.02531808],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.8896396,"threshold_uncertainty_score":0.5836484,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.0561997177844664,"score_gpt":0.38422726487057,"score_spread":0.3280275470861035,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}