{"id":"W2972918700","doi":"10.3390/pr7090614","title":"A Comparison of Clustering and Prediction Methods for Identifying Key Chemical–Biological Features Affecting Bioreactor Performance","year":2019,"lang":"en","type":"article","venue":"Processes","topic":"Spectroscopy and Chemometric Analyses","field":"Chemistry","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"University of British Columbia","funders":"","keywords":"Cluster analysis; Hierarchical clustering; Random forest; Computer science; Outcome (game theory); Data mining; Artificial intelligence; Machine learning; Feature selection; Biochemical engineering; Bioreactor; Set (abstract data type); Mathematics; Engineering; Biology","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.01832814,0.002028238,0.001441516,0.005702521,0.0006593987,0.001346244,0.001332578,0.001763534,0.0005615928],"category_scores_gemma":[0.02809907,0.0004418637,0.001787492,0.003165358,0.0006235767,0.001901796,0.0009355511,0.001364602,0.0003561056],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001201349,"about_ca_system_score_gemma":0.001275911,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005047667,"about_ca_topic_score_gemma":0.003688209,"domain_scores_codex":[0.9931341,0.003411892,0.0003833323,0.0008985481,0.001954482,0.0002175703],"domain_scores_gemma":[0.9670364,0.02600384,0.001390961,0.001700076,0.003606991,0.000261726],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"observational","study_design_scores_codex":[0.001570813,0.0005629579,0.03577818,0.001009274,0.00172053,0.0000676553,0.0003603878,0.280888,0.005062462,0.006044863,0.002832006,0.6641029],"study_design_scores_gemma":[0.00004585754,0.0006037911,0.01509581,0.0001490286,0.0002608048,0.0001117826,0.0001219125,0.9747174,0.003901063,0.003692144,0.001203998,0.00009642878],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.2237381,0.01191407,0.7553225,0.001488381,0.0002774022,0.0002776894,0.0006430856,0.00160452,0.004734321],"genre_scores_gemma":[0.7086105,0.002932339,0.2857926,0.0001953883,0.0001272483,0.0001962209,0.0009021696,0.0001788165,0.001064793],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01832814,"threshold_uncertainty_score":0.09692967,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.06544163702627569,"score_gpt":0.4064067179162489,"score_spread":0.3409650808899732,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}