{"id":"W2418468785","doi":"","title":"A method for verifying a vector-based text classification system.","year":2008,"lang":"en","type":"article","venue":"PubMed","topic":"Biomedical Text Mining and Ontologies","field":"Biochemistry, Genetics and Molecular Biology","cited_by":4,"is_retracted":false,"has_abstract":true,"ca_institutions":"Lockheed Martin (Canada)","funders":"","keywords":"Computer science; Search engine indexing; Java; Suite; Set (abstract data type); Lisp; Similarity (geometry); Data mining; Index (typography); Information retrieval; Vector space model; Artificial intelligence; Programming language","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.008642539,0.001147673,0.0007961098,0.006084044,0.001485099,0.003543256,0.001864087,0.001285044,0.01038995],"category_scores_gemma":[0.04844197,0.0005941307,0.001085015,0.004150104,0.001181898,0.004596066,0.002645463,0.001191121,0.00795547],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001138789,"about_ca_system_score_gemma":0.003359462,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003061643,"about_ca_topic_score_gemma":0.002395385,"domain_scores_codex":[0.9846297,0.002429212,0.003342747,0.002007463,0.007246392,0.0003444809],"domain_scores_gemma":[0.9755275,0.008453148,0.00189517,0.004364381,0.009372049,0.0003876509],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004819487,0.0002271416,0.007952257,0.001043447,0.0002001118,0.0003012175,0.0005925911,0.005698902,0.02785531,0.0450565,0.03019876,0.8803918],"study_design_scores_gemma":[0.0005025252,0.0009521453,0.01262423,0.000764819,0.0003391329,0.003619901,0.0009293362,0.4590199,0.1581993,0.1327528,0.2299294,0.0003665416],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.004549406,0.0001671017,0.9751873,0.000179447,0.0002314211,0.0004407211,0.001847598,0.01542295,0.001974025],"genre_scores_gemma":[0.04853808,0.0001233853,0.9429565,0.0001074949,0.00006183911,0.0007347587,0.004368493,0.0007013417,0.002408144],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01038995,"threshold_uncertainty_score":0.04570669,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.07073376003347205,"score_gpt":0.2900451600503665,"score_spread":0.2193114000168944,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}