{"id":"W2891172950","doi":"10.1007/978-3-030-00937-3_45","title":"How to Exploit Weaknesses in Biomedical Challenge Design and Organization","year":2018,"lang":"en","type":"book-chapter","venue":"Lecture notes in computer science","topic":"Biomedical and Engineering Education","field":"Engineering","cited_by":33,"is_retracted":false,"has_abstract":false,"ca_institutions":"McGill University","funders":"Deutsches Krebsforschungszentrum","keywords":"Benchmarking; Computer science; Exploit; Strengths and weaknesses; Data science; Best practice; Management science; Computer security; Political science; Management","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.005520464,0.0009173187,0.0005677114,0.0007751812,0.0009310406,0.003580789,0.001724228,0.001798325,0.01139457],"category_scores_gemma":[0.02213138,0.001028348,0.0007551191,0.0004473636,0.002169099,0.007377717,0.004919795,0.003189321,0.003898771],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007696493,"about_ca_system_score_gemma":0.001412867,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0009101792,"about_ca_topic_score_gemma":0.001361685,"domain_scores_codex":[0.9975613,0.0007020577,0.0001967951,0.0004095986,0.0009419393,0.0001882546],"domain_scores_gemma":[0.9899536,0.004209839,0.0005507092,0.00314281,0.001810986,0.0003318841],"domain_codex":null,"domain_gemma":"methods","domain_candidate":"methods","domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002849202,0.0001287232,0.003170965,0.0006603356,0.0001337965,0.0002386253,0.001374908,0.06484547,0.02928274,0.271594,0.0266136,0.6016719],"study_design_scores_gemma":[0.0000515178,0.0001586698,0.0006408404,0.0003222766,0.0001117068,0.0005208279,0.000746855,0.308237,0.03563851,0.537138,0.1163605,0.00007324622],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.008075923,0.0004279154,0.9738535,0.002676828,0.0001766536,0.0001030739,0.00009226019,0.001328036,0.01326577],"genre_scores_gemma":[0.1735638,0.0005173437,0.8050904,0.0008854123,0.00009993991,0.0002656477,0.0002511186,0.001031925,0.01829439],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.9944795,"threshold_uncertainty_score":0.0381186,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01552481616162634,"score_gpt":0.2065904991412381,"score_spread":0.1910656829796117,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}