{"id":"W4300167010","doi":"10.1016/j.jbiotec.2022.09.015","title":"Multiple genome analytics framework: The case of all SARS-CoV-2 complete variants","year":2022,"lang":"en","type":"article","venue":"Journal of Biotechnology","topic":"Genomics and Phylogenetic Studies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":3,"is_retracted":false,"has_abstract":false,"ca_institutions":"University of Calgary","funders":"","keywords":"Genome; Computer science; String (physics); Sequence (biology); Computational biology; Scalability; Computational complexity theory; Analytics; Pattern matching; Theoretical computer science; Data mining; Biology; Algorithm; Genetics; Artificial intelligence; Mathematics; Gene; Database","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004338201,0.001024522,0.0009377286,0.002340672,0.0009533037,0.005909015,0.002234432,0.001560897,0.004908782],"category_scores_gemma":[0.008879786,0.0005165137,0.001658856,0.002130119,0.001058424,0.004331856,0.003941371,0.002262052,0.00156127],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009545026,"about_ca_system_score_gemma":0.002304021,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007674824,"about_ca_topic_score_gemma":0.0105822,"domain_scores_codex":[0.9974731,0.0004870569,0.0002076986,0.0006078326,0.0009435867,0.0002808072],"domain_scores_gemma":[0.9967801,0.001225069,0.0002276596,0.0008520412,0.0006189239,0.0002961493],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002205043,0.0006104967,0.02649865,0.001314519,0.0005999312,0.007470803,0.003608042,0.1469423,0.03224826,0.3895124,0.05955376,0.3294358],"study_design_scores_gemma":[0.00008635549,0.0001530759,0.003314614,0.0002229329,0.0001747141,0.001185069,0.001126641,0.6059551,0.01424885,0.2787423,0.09466842,0.0001219582],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.04480639,0.001103669,0.8951465,0.003673039,0.0002247682,0.0002949683,0.007927394,0.03539539,0.01142792],"genre_scores_gemma":[0.2843347,0.000676046,0.6959118,0.0004730417,0.000121971,0.0001511632,0.01109721,0.00281854,0.004415688],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.007674824,"threshold_uncertainty_score":0.0229429,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.03335894066372589,"score_gpt":0.274429342463765,"score_spread":0.2410704018000392,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}