{"id":"W4405414691","doi":"10.1007/s00371-024-03729-0","title":"Dynamic text prompt joint multimodal features for accurate plant disease image captioning","year":2024,"lang":"en","type":"article","venue":"The Visual Computer","topic":"Multimodal Machine Learning Applications","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":false,"ca_institutions":"Carleton University","funders":"Natural Science Foundation of Hebei Province; National Natural Science Foundation of China","keywords":"Closed captioning; Computer science; Feature (linguistics); Artificial intelligence; Plant disease; Image (mathematics); Machine learning; Natural language processing","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005888276,0.001800788,0.0008140012,0.002032194,0.0004541577,0.001330549,0.0009944864,0.00181693,0.03285933],"category_scores_gemma":[0.003239055,0.0003545905,0.0007178858,0.0009693779,0.000248551,0.001568904,0.001327356,0.001355219,0.01886473],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004736701,"about_ca_system_score_gemma":0.0005706974,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001752679,"about_ca_topic_score_gemma":0.002654827,"domain_scores_codex":[0.9995986,0.00008127545,0.00002491754,0.0001075975,0.0001112336,0.00007648991],"domain_scores_gemma":[0.9985992,0.0003624301,0.00009468053,0.0001931851,0.0006551518,0.00009532279],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001927049,0.0002476524,0.0007036415,0.0009102911,0.00005175975,0.0005060548,0.00009892252,0.007976475,0.1498351,0.001323376,0.1058119,0.7306079],"study_design_scores_gemma":[0.0001249682,0.0005543262,0.006056574,0.0002487732,0.0001660445,0.0008224538,0.0002837995,0.6934956,0.2100859,0.005088169,0.08294562,0.0001277093],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.08264563,0.005525484,0.749056,0.002111324,0.004259675,0.001233595,0.0268304,0.09961146,0.02872648],"genre_scores_gemma":[0.3682091,0.002610515,0.5610369,0.001214665,0.001543034,0.001082252,0.03031878,0.004216117,0.02976862],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.03285933,"threshold_uncertainty_score":0.1099254,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01490606545697368,"score_gpt":0.3184719595735442,"score_spread":0.3035658941165705,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}