<?xml version="1.0" encoding="UTF-8" standalone="no"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.3 20070202//EN" "journalpublishing.dtd">
<article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" article-type="research-article" dtd-version="2.3" xml:lang="EN">
<front>
<journal-meta>
<journal-id journal-id-type="publisher-id">Front. Plant Sci.</journal-id>
<journal-title>Frontiers in Plant Science</journal-title>
<abbrev-journal-title abbrev-type="pubmed">Front. Plant Sci.</abbrev-journal-title>
<issn pub-type="epub">1664-462X</issn>
<publisher>
<publisher-name>Frontiers Media S.A.</publisher-name>
</publisher>
</journal-meta>
<article-meta>
<article-id pub-id-type="doi">10.3389/fpls.2025.1612800</article-id>
<article-categories>
<subj-group subj-group-type="heading">
<subject>Plant Science</subject>
<subj-group>
<subject>Original Research</subject>
</subj-group>
</subj-group>
</article-categories>
<title-group>
<article-title>Developing sustainable system based on transformers algorithms to predict the Dubas insects diseases in palm leaves</article-title>
</title-group>
<contrib-group>
<contrib contrib-type="author" corresp="yes">
<name>
<surname>Aldhyani</surname>
<given-names>Theyazn H. H.</given-names>
</name>
<xref ref-type="aff" rid="aff1">
<sup>1</sup>
</xref>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<xref ref-type="author-notes" rid="fn001">
<sup>*</sup>
</xref>
<uri xlink:href="https://loop.frontiersin.org/people/1076026/overview"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/visualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/conceptualization/"/>
<role content-type="https://credit.niso.org/contributor-roles/data-curation/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
</contrib>
<contrib contrib-type="author">
<name>
<surname>Alkahtani</surname>
<given-names>Hasan</given-names>
</name>
<xref ref-type="aff" rid="aff2">
<sup>2</sup>
</xref>
<role content-type="https://credit.niso.org/contributor-roles/writing-review-editing/"/>
<role content-type="https://credit.niso.org/contributor-roles/software/"/>
<role content-type="https://credit.niso.org/contributor-roles/validation/"/>
<role content-type="https://credit.niso.org/contributor-roles/funding-acquisition/"/>
<role content-type="https://credit.niso.org/contributor-roles/resources/"/>
<role content-type="https://credit.niso.org/contributor-roles/formal-analysis/"/>
<role content-type="https://credit.niso.org/contributor-roles/writing-original-draft/"/>
<role content-type="https://credit.niso.org/contributor-roles/methodology/"/>
<role content-type="https://credit.niso.org/contributor-roles/supervision/"/>
<role content-type="https://credit.niso.org/contributor-roles/investigation/"/>
<role content-type="https://credit.niso.org/contributor-roles/project-administration/"/>
</contrib>
</contrib-group>
<aff id="aff1">
<sup>1</sup>
<institution>Applied College, King Faisal University</institution>, <addr-line>Al-Ahsa</addr-line>,&#xa0;<country>Saudi Arabia</country>
</aff>
<aff id="aff2">
<sup>2</sup>
<institution>College of Computer Science and Information Technology, King Faisal University</institution>, <addr-line>Al-Ahsa</addr-line>,&#xa0;<country>Saudi Arabia</country>
</aff>
<author-notes>
<fn fn-type="edited-by">
<p>Edited by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/788131/overview">Anirban Roy</ext-link>, Indian Council of Agricultural Research (ICAR), India</p>
</fn>
<fn fn-type="edited-by">
<p>Reviewed by: <ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1484190/overview">Preeta Sharan</ext-link>, The Oxford College of Engineering, India</p>
<p>
<ext-link ext-link-type="uri" xlink:href="https://loop.frontiersin.org/people/1763479/overview">Aibin Chen</ext-link>, Central South University Forestry and Technology, China</p>
</fn>
<fn fn-type="corresp" id="fn001">
<p>*Correspondence: Theyazn H. H. Aldhyani, <email xlink:href="mailto:taldhyani@kfu.edu.sa">taldhyani@kfu.edu.sa</email>
</p>
</fn>
</author-notes>
<pub-date pub-type="epub">
<day>04</day>
<month>09</month>
<year>2025</year>
</pub-date>
<pub-date pub-type="collection">
<year>2025</year>
</pub-date>
<volume>16</volume>
<elocation-id>1612800</elocation-id>
<history>
<date date-type="received">
<day>16</day>
<month>04</month>
<year>2025</year>
</date>
<date date-type="accepted">
<day>12</day>
<month>08</month>
<year>2025</year>
</date>
</history>
<permissions>
<copyright-statement>Copyright &#xa9; 2025 Aldhyani and Alkahtani.</copyright-statement>
<copyright-year>2025</copyright-year>
<copyright-holder>Aldhyani and Alkahtani</copyright-holder>
<license xlink:href="http://creativecommons.org/licenses/by/4.0/">
<p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (CC BY). The use, distribution or reproduction in other forums is permitted, provided the original author(s) and the copyright owner(s) are credited and that the original publication in this journal is cited, in accordance with accepted academic practice. No use, distribution or reproduction is permitted which does not comply with these terms.</p>
</license>
</permissions>
<abstract>
<sec>
<title>Introduction</title>
<p>Agriculture has emerged as a crucial area of inquiry, presenting a significant challenge for numerous experts in the field of computer vision. Identifying and categorizing plant diseases at an early stage is essential for mitigating the spread of these diseases and preventing a decline in crop yields. The overall condition of palm trees, including their roots, stems, and leaves, plays a crucial role in palm production, necessitating careful observation to ensure maximum yield. A significant challenge in maintaining productive crops is the widespread presence of pests and diseases that affect palm plants. The impact of these diseases on growth and development can be significantly negative, resulting in reduced productivity. The productivity of palms is intricately linked to the state of their leaves, which are essential for the process of photosynthesis.</p>
</sec>
<sec>
<title>Methods</title>
<p>This study utilized an extensive dataset comprising 1600 images, which included 800 images of healthy leaves and another 800 of Dubas images. Additionally, the primary aim was to develop EfficientNetV2B0, DenseNet12, and a transformer model known as the Vision Transformer (ViT) model for detecting diseases and pests affecting palm leaves, utilizing image analysis methods to enhance pest management strategies.</p>
</sec>
<sec>
<title>Results</title>
<p>The proposed models demonstrated superior performance compared to numerous recent studies in the field, utilizing established metrics on both original and augmented datasets, achieving an impressive accuracy of 99.37% with the ViT model.</p>
</sec>
<sec>
<title>Discussion</title>
<p>This study presents an innovative approach for identifying diseases in palm leaves. This will have a significant impact on the agricultural sector. The results were quite promising, justifying their implementation in palm companies to improve pest and disease management</p>
</sec>
</abstract>
<kwd-group>
<kwd>palm</kwd>
<kwd>diseases</kwd>
<kwd>transformers</kwd>
<kwd>deep learning</kwd>
<kwd>sustainable</kwd>
<kwd>insect</kwd>
</kwd-group>
<counts>
<fig-count count="14"/>
<table-count count="6"/>
<equation-count count="0"/>
<ref-count count="37"/>
<page-count count="15"/>
<word-count count="5463"/>
</counts>
<custom-meta-wrap>
<custom-meta>
<meta-name>section-in-acceptance</meta-name>
<meta-value>Sustainable and Intelligent Phytoprotection</meta-value>
</custom-meta>
</custom-meta-wrap>
</article-meta>
</front>
<body>
<sec id="s1" sec-type="intro">
<label>1</label>
<title>Introduction</title>
<p>Round 70% of the world&#x2019;s date fruit comes from Saudi Arabia, thanks to its more than thirty-one million palm trees (<xref ref-type="bibr" rid="B31">The Ministry of Environment and Water and Agriculture Saudi Arabia, 2020</xref>). In 2023, the more than 26,000 date farms located in the Madinah area of western Saudi Arabia produced about $253 million. Roots, trunks, leaves, and fruits of palm trees are all vulnerable to a host of infectious illnesses caused by bacteria and fungi (<xref ref-type="bibr" rid="B32">The Saudi Press Agency (SPA), 2024</xref>).</p>
<p>The inventory of palm trees is essential for assessing diversity and health; however, data regarding their numbers and distribution is limited, outdated, and inconsistent. Estimations of date palm trees in plantations rely on assessments that exclude non-agricultural areas and natural populations. Data on canary palm populations is restricted to areas of significant public interest, with estimates derived from date palm production rather than geospatial databases (<xref ref-type="bibr" rid="B16">Jaradat, 2015</xref>; <xref ref-type="bibr" rid="B5">Sharma et al., 2021</xref>; <xref ref-type="bibr" rid="B36">Zaid and de Wet, 2002</xref>).</p>
<p>Agriculture, and date palm trees in particular, are very vulnerable to diseases and climate change. Date palm trees suffer greatly from serious diseases, including Brittle Leaf disease, Brown Leaf Spot, Bayoud disease, Black Scorch, SDS, and Brown Leaf Spot, which drastically reduce fruit quality and productivity. This study examines the detection of the SDS disease, which disseminates in a hazardous manner, complicating control efforts and resulting in substantial losses in fruit output. This disease poses a significant threat to Date palm farming worldwide and hinders fresh planting efforts (<xref ref-type="bibr" rid="B18">Khamparia et&#xa0;al., 2020</xref>).</p>
<p>Palm leaves are integral to numerous ecosystems, economies, and civilizations. They play a crucial role in various aspects, supporting individuals in earning a livelihood and staying on their chosen course in life. Nevertheless, the Dubas bug presents a significant risk to date palm trees and their leaves, especially in areas where date palm farming is prevalent. Identifying and classifying plant diseases is crucial for precision agriculture; however, farmers face challenges in diagnosing these diseases and assessing the extent of infestation. The combination of machine vision and deep learning has revolutionized the automated detection and assessment of pests and diseases in agriculture. Palm leaf diseases pose a considerable challenge to the health and productivity of trees (<xref ref-type="bibr" rid="B25">Resh and Card&#xe9;, 2009</xref>; <xref ref-type="bibr" rid="B14">Hessane et&#xa0;al., 2023</xref>).</p>
<p>Adult, nymph, and egg are the three phases that make up a dubas bug&#x2019;s life cycle. With its two sets of wings, this insect is hemimetabolous. It undergoes a mating cycle in the spring (February&#x2013;May) and another in the fall (August&#x2013;November) of each year. Date palm females use the third and fifth leaf fronds, respectively, to deposit their eggs in the spring and autumn.</p>
<p>The life cycle begins with oviposition, followed by hatching into nymphs and undergoing five molts until the adult form is attained. Adults have a yellowish-brown to greenish coloration, characterized by two black patches on their heads. They generate honeydew and necrotic regions in plant tissues due to their oviposition behavior. Nonetheless, it remains uncertain whether these necrotic lesions result from fungal infections (<xref ref-type="bibr" rid="B10">Elwan and Al-Tamimi, 1999</xref>; <xref ref-type="bibr" rid="B30">Thacker et&#xa0;al., 2003</xref>; <xref ref-type="bibr" rid="B8">Ba-Angood et&#xa0;al., 2009</xref>; <xref ref-type="bibr" rid="B25">Resh and Card&#xe9;, 2009</xref>; <xref ref-type="bibr" rid="B24">Payandeh et&#xa0;al., 2010</xref>).</p>
<p>Computer-based technologies are rapidly revolutionizing agriculture, reducing human labor, and enabling impartial decision-making. Image processing methods are utilized in various computer vision applications for disease diagnosis, identification, and segmentation tasks. This technological development is revolutionizing agriculture, improving efficiency and efficacy. Using a modified MobileNetV2 neural network, the authors (<xref ref-type="bibr" rid="B19">Kong et&#xa0;al., 2022</xref>) enhanced the accuracy of cassava leaf disease classification by employing data augmentation methods on lower-quality test images.</p>
<p>With high-quality photographs, they achieved a 97% recognition accuracy; however, this accuracy was significantly reduced with low-quality images. Using a range of classifiers at the image level, including Fine KNN, Cubic SVM, and tree ensemble, the authors in (<xref ref-type="bibr" rid="B20">Li et&#xa0;al., 2023</xref>) classified guava plant illnesses with an overall classification accuracy of 99%. Plant disease classification is achieved by a hybrid wrapper model that combines CNN classifiers with FPA-SVM (<xref ref-type="bibr" rid="B9">Bayomi-Alli et&#xa0;al., 2021</xref>). This methodology yielded a classifier with an accuracy of 99% through feature selection using FPA and SVM in a wrapper approach. The authors (<xref ref-type="bibr" rid="B2">Almadhor et&#xa0;al., 2021</xref>) proposed a deep learning (DL) model for disease detection on cucumber and potato leaves, utilizing an optimization approach. A 99% success rate was achieved by optimizing the DFs obtained from the global pooling layer using an upgraded Cuckoo search strategy. For disease categorization in plant leaves, <xref ref-type="bibr" rid="B35">Ya&#x11f; and Altan (2022)</xref> presented the EfficientNet DL architecture. They employed a transfer learning approach to train their model and other deep learning models, and both performed well. To classify citrus diseases, <xref ref-type="bibr" rid="B37">Zia Ur Rehman et&#xa0;al. (2022)</xref> developed a new DL model. The accuracy rate was 95% because the Whale Optimisation Algorithm (WOA) was used to retrain two pre-trained models, DenseNet 201 and MobileNetv2, so that they could produce feature vectors. A DL model for guava disease identification has achieved a 97% accuracy rate by utilizing enhanced data supplemented with color-histogram equalization and unsharp masking techniques (<xref ref-type="bibr" rid="B7">Atila et&#xa0;al., 2021</xref>). The primary contributions of this research are as follows.</p>
<p>Make a substantial contribution to the categorization of palm leaf diseases by using cutting-edge deep learning architecture. This paper presents a method for the automated detection and enumeration of palm leaves, disease identification, and assessment of palm health from high-resolution images using deep learning and Vision Transformer models. The training dataset included more than 1600 images of individual dubas and healthy classifications. We assessed the model via training assessment and by comparing prediction outcomes with visual and ground inspections. The model was also evaluated using images captured at varying elevations. It can achieve elevated accuracy with less labeled data by using pre-trained models. This method enhances classification efficacy while reducing training duration and computational expenses. Additionally, it streamlines the process of disease identification, providing a scalable and rapid option for early detection. This is essential for agricultural disease management.</p>
</sec>
<sec id="s2">
<label>2</label>
<title>Related works</title>
<p>Several new approaches to identifying and categorizing plant diseases have recently been developed. The methods were tested on a variety of datasets, each with its distinct features. Representing plant photos using practical and discriminative characteristics is essential for creating a system to identify plant diseases. There are two primary schools of thought regarding feature extraction methods: those that rely on manually created features and those that utilize deep learning techniques (<xref ref-type="bibr" rid="B28">Shah et&#xa0;al., 2016</xref>; <xref ref-type="bibr" rid="B27">Sethy et&#xa0;al., 2020</xref>; <xref ref-type="bibr" rid="B15">Iqbal et&#xa0;al., 2023</xref>; <xref ref-type="bibr" rid="B33">Tholkapiyan et&#xa0;al., 2023</xref>).</p>
<p>Using CNN with SVM, <xref ref-type="bibr" rid="B17">Kamal et&#xa0;al. (2018)</xref>. We were able to differentiate between Chimara (the most prevalent date palm spot leaf disease) and Anthracnose (the least frequent), achieving accuracy rates of 97% and 95%, respectively. To attain a classification accuracy of 99.67% with an artificial neural network (ANN) classifier, <xref ref-type="bibr" rid="B12">Hamdani et&#xa0;al. (2021)</xref> suggested an approach that makes use of a color histogram feature and a dataset consisting of 300 laboratory photos. With an overall accuracy of above 96%, <xref ref-type="bibr" rid="B21">Liu et&#xa0;al. (2021)</xref> used a DL-based Faster RCNN to identify and quantify oil palm plants in UAV pictures. <xref ref-type="bibr" rid="B1">Abu-zanona et&#xa0;al. (2022)</xref> achieved the best accuracy for the Kaggle dataset, classifying four types of sick palm trees with 97% accuracy using VGG 16 and MobileNet.</p>
<p>
<xref ref-type="bibr" rid="B26">Septiarini et&#xa0;al. (2021)</xref> presented a method for detecting diseases in oil palm leaves. Otsu thresholding was employed in the Lab color space to identify Regions of Interest (ROIs), followed by preprocessing and classification using k-nearest neighbors (KNN) (<xref ref-type="bibr" rid="B11">Eunice et&#xa0;al., 2022</xref>). The identification of leaf diseases in tomato plants was conducted by Sunil S. Harakannanavar et&#xa0;al (<xref ref-type="bibr" rid="B13">Harakannanavar et&#xa0;al., 2022</xref>). This approach combines multiple techniques, such as SVM, KNN, and CNN, for the detection of palm diseases. Recent studies have developed a CNN for disease classification in palm trees (<xref ref-type="bibr" rid="B1">Abu-zanona et&#xa0;al., 2022</xref>).</p>
<p>Authors (<xref ref-type="bibr" rid="B23">Nusrat et&#xa0;al., 2020</xref>) were able to detect and categorize wheat illness leaves with a 98% success rate using SVM and GoogLeNet, two ML models. A prior study by (<xref ref-type="bibr" rid="B34">Too et&#xa0;al., 2019</xref>) compared several DL models for plant disease detection. This study utilized 121-layer customized models, including VGG-16, Inception V4, ResNet, and DenseNets, to classify plant species. This study used the model to detect and classify healthy leaves, as well as four prevalent diseases that may affect palm trees: bacterial leaf blight, brown spots, leaf smut, and white scale. When tested against VGG-16 and MobileNet, two prominent CNN models, the suggested model outperformed them both with an accuracy rate exceeding 99%. In a separate study, <xref ref-type="bibr" rid="B12">Hamdani et&#xa0;al. (2021)</xref> examined methods for identifying and categorizing diseases that affect palm trees. The authors of this study achieved a 99% success rate in classifying palm diseases using an artificial neural network (ANN) classifier and principal component analysis (PCA) to extract color features.</p>
<p>
<xref ref-type="bibr" rid="B22">Masazhar and Kamal (2017)</xref> employed an automated system to detect and categorize disease indicators in palm oil leaves. Two palm oil illnesses were successfully detected using k-means clustering, an SVM classifier, and leaf symptomatology. Thirteen novels were produced from k-means clustering for disease classification. <xref ref-type="bibr" rid="B6">Ashqar and Abu-Naser (2018)</xref> utilized a CNN model for identifying tomato leaf diseases, demonstrating improved performance with full-color images compared to grayscale images. Shruthi U et&#xa0;al. highlighted the efficacy of convolutional neural networks in identifying specific agricultural diseases through machine learning methodologies (<xref ref-type="bibr" rid="B29">Shruthi et&#xa0;al., 2019</xref>).</p>
</sec>
<sec id="s3" sec-type="materials|methods">
<label>3</label>
<title>Materials and methods</title>
<sec id="s3_1">
<label>3.1</label>
<title>Farmwork of the proposed system</title>
<p>The suggested paradigm for disease detection in palm leaves is shown in <xref ref-type="fig" rid="f1">
<bold>Figure&#xa0;1</bold>
</xref>. The proposed framework covers the following crucial steps. The data is first enhanced to improve the training process by increasing the sample count. Then, we choose and improve EfficientNetV2B0, DenseNet12, and Vision Transformer. ViT features are obtained from the global pooling layer, and deep learning further trains the models. We conclude by drawing parallels to popular and up-to-date DL and transformer kinds. Training Vision Transform models using the Palm-leaves dataset is the focus of this research.</p>
<fig id="f1" position="float">
<label>Figure&#xa0;1</label>
<caption>
<p>Farmwork of the palm system.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g001.tif">
<alt-text content-type="machine-generated">Flowchart detailing a machine learning process for classifying leaf images as either &#x201c;Dubas&#x201d; or &#x201c;Healthy.&#x201d; It begins with a dataset, involves importing libraries, data splitting, and augmentation. Two models are depicted: a deep learning model with convolution and pooling layers, and a Vision Transformer model with attention mechanisms. Both models lead to output layers, followed by training, fine-tuning, and evaluation, resulting in predictions of leaf health.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_2">
<label>3.2</label>
<title>Dataset</title>
<p>We obtained the dataset from the Karbala Governorate in Iraq via Kaggle, collecting leaves to varying degrees. In the research, we used palm leaves with images of Dubas and health. The image resolutions are 6000 &#xd7; 4000 &#xd7; 3 pixels for the Canon 77D camera and 8000 &#xd7; 6000 &#xd7; 3 pixels for the DJI Camera 800 images of Dubas and 800 images of the health class. A snapshot of the palm leaves is presented in <xref ref-type="fig" rid="f2">
<bold>Figure&#xa0;2</bold>
</xref>. <xref ref-type="fig" rid="f3">
<bold>Figure&#xa0;3</bold>
</xref> presents the class values of the palm leaves dataset.</p>
<fig id="f2" position="float">
<label>Figure&#xa0;2</label>
<caption>
<p>Sample from palm data.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g002.tif">
<alt-text content-type="machine-generated">Upper row shows close-up images of leaves with &#x201c;Dubas&#x201d; infestation, featuring visible brown spots and blemishes. Lower row shows &#x201c;Healthy&#x201d; leaves with smooth, unblemished green surfaces.</alt-text>
</graphic>
</fig>
<fig id="f3" position="float">
<label>Figure&#xa0;3</label>
<caption>
<p>Class of palm data.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g003.tif">
<alt-text content-type="machine-generated">Bar chart showing class distribution with two categories: &#x201c;Dubas&#x201d; and &#x201c;Healthy.&#x201d; Both categories have around 800 images. &#x201c;Dubas&#x201d; is represented in blue and &#x201c;Healthy&#x201d; in green.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_3">
<label>3.2</label>
<title>Preprocessing steps</title>
<p>During the preprocessing phase of EfficientNetV2B0, DenseNet121, and ViT, several critical processing steps were implemented to ensure high-quality input data for model training. Initially, all images were downsized to a consistent dimension suitable for each model&#x2019;s architecture: 224&#xd7;224&#xd7;3 for EfficientNetV2B0 and DenseNet121, and 384&#xd7;384&#xd7;3 for ViT. Subsequently, the pixel values of palm images were standardized using the mean and standard deviation to enhance model convergence. These models used data augmentation methods, including rotation, flipping, zooming, and contrast modification, to improve model generalization and flexibility. Furthermore, the ViT model uses the tokenization of image patches before processing, whereas CNN-based models employ feature scaling to ensure consistency. These preprocessing measures enhance model efficiency, accuracy, and generalizability.</p>
</sec>
<sec id="s3_4">
<label>3.3</label>
<title>Proposed models</title>
<sec id="s3_4_1">
<label>3.3.1</label>
<title>EfficientNet-B0 model</title>
<p>The EfficientNet-B0 architecture is a well-established CNN model that can serve as an encoder in tasks involving semantic segmentation. EfficientNet-B0 served as the backbone network in the proposed research design for feature extraction from the input image through downsampling. EfficientNet-B0 is a CNN architecture comprising several blocks, each incorporating convolutional layers, activation functions, and pooling operations. This architecture is a convolutional neural network commonly employed for image classification tasks. The output of EfficientNet-B0 is frequently used as input for a decoder network in semantic segmentation. The application of EfficientNet-B0 as an encoder for semantic segmentation has demonstrated remarkable accuracy and efficiency in various contexts, particularly in medical image segmentation (<xref ref-type="bibr" rid="B12">Hamdani et&#xa0;al., 2021</xref>; <xref ref-type="bibr" rid="B21">Liu et&#xa0;al., 2021</xref>). <xref ref-type="fig" rid="f4">
<bold>Figure&#xa0;4</bold>
</xref> presents the plots generated by the encoder model.</p>
<fig id="f4" position="float">
<label>Figure&#xa0;4</label>
<caption>
<p>EfficientNetV2B0 model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g004.tif">
<alt-text content-type="machine-generated">Flowchart illustrating a deep learning model architecture. The process begins with an input of size 224x224x3, followed by a convolutional layer with 32 filters, stride of 2. It includes fused-MBConv blocks with varying expansions and channels, progressing through batch normalization and SiLU activation. The architecture then transitions to MBConv6 blocks with specific expansions and SE ratios, culminating in a global average pooling and output layer.</alt-text>
</graphic>
</fig>
<p>The network is optimized for classification (healthy vs. Dubas) through the incorporation of a global average pooling layer and a dense layer utilizing a sigmoid activation function. The model was trained on augmented image data that had horizontal flips. An adaptive learning rate and early stopping were used to avoid overfitting. Using depthwise separable convolutions, batch normalization, and activation layers, the EfficientNetV2B0 architecture has 236 layers that are optimized for effective feature extraction. A transfer learning framework utilizes the pre-trained base, omitting the upper classification layers. Along with the base model, two extra layers are added: a Global Average Pooling layer that combines spatial data from feature maps and a Dense layer that uses a sigmoid activation function for binary classification. The model comprises 238 layers and combines EfficientNetV2B0&#x2019;s powerful feature-extraction capabilities with a simplified custom classification head designed for binary classification tasks. <xref ref-type="table" rid="T1">
<bold>Table&#xa0;1</bold>
</xref> displays the parameters of the EfficientNet-B0 model.</p>
<table-wrap id="T1" position="float">
<label>Table&#xa0;1</label>
<caption>
<p>Parameter EfficientNet-B0 model.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">#Name</th>
<th valign="middle" align="center">#Values</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Layers</td>
<td valign="middle" align="center">238</td>
</tr>
<tr>
<td valign="middle" align="center">Image</td>
<td valign="middle" align="center">224x224x3</td>
</tr>
<tr>
<td valign="middle" align="center">Optimize</td>
<td valign="middle" align="center">Adam</td>
</tr>
<tr>
<td valign="middle" align="center">Learning_rate</td>
<td valign="middle" align="center">0.001</td>
</tr>
<tr>
<td valign="middle" align="center">Batch_Size</td>
<td valign="middle" align="center">16</td>
</tr>
<tr>
<td valign="middle" align="center">Epochs</td>
<td valign="middle" align="center">20</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
<sec id="s3_4_2">
<label>3.3.2</label>
<title>DenseNet121 model</title>
<p>DenseNet121 is a CNN architecture proposed by Huang et&#xa0;al. It belongs to the DenseNet family, recognized for its dense network architecture and remarkable efficacy in several computer vision applications, including image categorization. The design of DenseNet-121, as shown in <xref ref-type="fig" rid="f5">
<bold>Figure&#xa0;5</bold>
</xref>, is centered on the principle of dense connections. Unlike standard CNN designs, where layers are sequentially linked, DenseNet utilizes skip connections that link each layer to every other layer in a feed-forward way. This intricate connection structure facilitates direct feature reuse and promotes information flow throughout the network, leading to improved gradient propagation, enhanced feature extraction, and overall model efficacy. The core feature extractor is DenseNet-121, a pre-trained convolutional neural network developed for image classification tasks. Across all 121 levels, this design effectively utilizes gradient movement and feature reuse due to its rich connections. Class weights, a binary cross-entropy loss function, and the Adam optimizer, which has a learning rate of 0.001, are used by the model to handle data imbalances. Metrics, including confusion matrices, classification reports, and accuracy, are used for evaluation. Essential parameters are a batch size of 16, a target image dimension of (224, 224), and 20 training epochs.</p>
<fig id="f5" position="float">
<label>Figure&#xa0;5</label>
<caption>
<p>DenseNet121 model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g005.tif">
<alt-text content-type="machine-generated">Flowchart of a neural network model, likely a DenseNet, showing input processing with convolution and pooling, followed by four dense blocks with transition layers. Each block consists of multiple layers: Dense Block 1 has 6 layers, Dense Block 2 has 12 layers, Dense Block 3 has 24 layers, and Dense Block 4 has 16 layers. Transition layers have channels of 256, 512, and 1024 respectively. After Block 4, global average pooling reduces features to 1024 dimensions, followed by a 1000-D features layer. Arrows indicate processing flow.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s3_4_3">
<label>3.3.3</label>
<title>Vision transformer model</title>
<p>The Vision Transformer (ViT) is an innovative neural network design that transforms the processing and comprehension of images. The Vision Transformer (ViT) concept was presented in 2021 in a conference research paper entitled &#x201c;An Image is Worth 16*16 Words.&#x201d; Transformers for Image Recognition at Scale, or ViT, presents an innovative approach to image analysis by segmenting images into smaller patches and using self-attention processes. This enables the model to discern both local and global links among images, resulting in remarkable performance across many computer vision tasks. Whereas CNNs immediately analyze raw pixel values, ViT segments the input image into patches and converts them into tokens. ViT utilizes self-attention processes to analyze connections among all patches. The Vision Transformer (ViT) inherently captures global context through self-attention, enabling the recognition of relationships among distant patches. Convolutional Neural Networks use pooling layers to extract coarse global information.</p>
<p>It can be fine-tuned for individual tasks after pre-training on large datasets. By segmenting 224x224 input images into patches and modeling their interrelationships using self-attention mechanisms, the ViT model can perform analysis. The B16 architecture is used by this ViT model, as seen in <xref ref-type="fig" rid="f6">
<bold>Figure&#xa0;6</bold>
</xref>. Feature extraction is made easier with its transformer block, which incorporates feedforward neural networks and many attention layers trained on massive datasets. An integrated custom classification head, along with immobilized pre-trained weights and layers, ensures the model&#x2019;s obtained representations remain intact. This head comprises an activation layer with a regularization rate of 0.5, a dense layer with 128 neurons activated by ReLU, and a final dense layer with sigmoid activation for binary classification. Class weights are used to reduce class imbalance, and adjustments to data augmentation methods improve the model&#x2019;s generalizability. After 20 iterations of training using the Adam optimizer at a learning rate of 0.001, the model is complete. The ViT values are displayed in <xref ref-type="table" rid="T2">
<bold>Table&#xa0;2</bold>
</xref>.</p>
<fig id="f6" position="float">
<label>Figure&#xa0;6</label>
<caption>
<p>ViT model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g006.tif">
<alt-text content-type="machine-generated">Diagram of a neural network architecture with processes labeled from input to output. The flow includes multi-head attention, add and normalize layers, feed forward, and finally a flatten layer before output. Arrows indicate data flow direction.</alt-text>
</graphic>
</fig>
<table-wrap id="T2" position="float">
<label>Table&#xa0;2</label>
<caption>
<p>ViT parameters.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="left">#Name</th>
<th valign="middle" align="left">#Values</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="left">Patch_size:</td>
<td valign="middle" align="left">16,</td>
</tr>
<tr>
<td valign="middle" align="left">Transformer _encoder:</td>
<td valign="middle" align="left">12</td>
</tr>
<tr>
<td valign="middle" align="left">Attention_heads</td>
<td valign="middle" align="left">12</td>
</tr>
<tr>
<td valign="middle" align="left">Hidden_size</td>
<td valign="middle" align="left">768</td>
</tr>
<tr>
<td valign="middle" align="left">Optimizer:</td>
<td valign="middle" align="left">Adam</td>
</tr>
<tr>
<td valign="middle" align="left">Epochs:</td>
<td valign="middle" align="left">20</td>
</tr>
<tr>
<td valign="middle" align="left">Learning_Rate:</td>
<td valign="middle" align="left">0.001</td>
</tr>
<tr>
<td valign="middle" align="left">Image_Size</td>
<td valign="middle" align="left">244X244</td>
</tr>
<tr>
<td valign="middle" align="left">Dropout_Rate:</td>
<td valign="middle" align="left">0.5</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
</sec>
<sec id="s4">
<label>4</label>
<title>Experiments and results discussion</title>
<p>The experiments were evaluated using a GPU P100 Kaggle system as the baseline. Operating Windows 11, the machine features 16 GB of RAM and a 9th-generation Core i7 CPU. Aiming to enable deep learning applications with reduced memory usage and improved execution speed, the software implementation included libraries such as Anaconda, Keras, OpenCV, NumPy, and CuDNN. For every experiment carried out, this work assessed the training and testing accuracy. Every model has calculated losses throughout the testing and training periods. Training the models on the Palm tree dataset helped to improve the learning speed of the transformer and transfer learning models. EfficientNetV2B0, DenseNet121, and ViT models were used for this work.</p>
<p>Each of the two dataset classes corresponded to a different disease. Since the Palm dataset&#x2019;s color images worked well with the DL and ViT models, we used them in our experiments. The images were downscaled to a uniform pixel format since different pre-trained network models require varied input sizes. Input dimensions of 224 &#xd7; 224 &#xd7; 3 (height, width, and channel depth) are used by EfficientNet V2B0, DenseNet 121, and ViT. Data augmentation after preprocessing is a regularization strategy that is used to reduce the impact of overfitting. This method of model augmentation makes the model more robust, which in turn enhances its ability to categorize images of real plant diseases while reducing the likelihood of overfitting and model loss.</p>
<sec id="s4_1" sec-type="results">
<label>4.1</label>
<title>Results</title>
<sec id="s4_1_1">
<label>4.1.1</label>
<title>Result of EfficientNetV2B0 model</title>
<p>The findings of the EfficientNetV2B0 model, presented in <xref ref-type="table" rid="T3">
<bold>Table&#xa0;3</bold>
</xref>, demonstrate robust performance across all assessed criteria, indicating a highly efficient classification system. The model demonstrates balanced and consistent performance in accurately detecting instances of both the &#x201c;Dubas&#x201d; and &#x201c;Healthy&#x201d; classes, with accuracy, recall, and F1-score all at 98%. The total accuracy of 98% further substantiates the model&#x2019;s reliability. The weighted average of accuracy, recall, and F1-score, again at 98%, indicates that the model generalizes well across the dataset, maintaining excellent performance without significant bias towards any one class. The findings underscore the efficacy of the EfficientNetV2B0 model for this classification job, demonstrating exceptional predicted accuracy and dependability.</p>
<table-wrap id="T3" position="float">
<label>Table&#xa0;3</label>
<caption>
<p>Performance of EfficientNetV2B0 model.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Class name</th>
<th valign="middle" align="center">Precision (%)</th>
<th valign="middle" align="center">Recall (%)</th>
<th valign="middle" align="center">F1-score (%)</th>
<th valign="middle" align="center">Support samples of/validation phase</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">96</td>
<td valign="middle" align="center">100</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">80</td>
</tr>
<tr>
<td valign="middle" align="center">Healthy</td>
<td valign="middle" align="center">100</td>
<td valign="middle" align="center">96</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">80</td>
</tr>
<tr>
<td valign="middle" align="center">Accuracy</td>
<td valign="middle" align="center"/>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center"/>
<td valign="middle" align="center"/>
</tr>
<tr>
<td valign="middle" align="center">Weighted_Avg_ palm system</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">160</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>
<xref ref-type="fig" rid="f7">
<bold>Figure&#xa0;7</bold>
</xref> illustrates the confusion matrix for the EfficientNetV2B0 model used in a classification job differentiating between healthy and Dubas samples, demonstrating outstanding performance. The model analyzed 160 test samples (80 each class), accurately identifying 80 healthy and 77 Dubas samples, resulting in 3 misclassifications from each class, yielding 3 false positives and 0 false negatives. This yields an exceptional accuracy of 98.12%, demonstrating the model&#x2019;s robust discriminative capability. Nonetheless, the almost flawless classification requires further validation on an independent dataset to verify its robustness and alleviate concerns over possible overfitting.</p>
<fig id="f7" position="float">
<label>Figure&#xa0;7</label>
<caption>
<p>confusion matrix EfficientNetV2B0 model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g007.tif">
<alt-text content-type="machine-generated">Confusion matrix titled &#x201c;Test Set Confusion Matrix&#x201d; with predicted vs. actual values. Top left: 80 true negatives; top right: 0 false positives; bottom left: 3 false negatives; bottom right: 77 true positives. A color gradient indicates values, with a dark blue for higher numbers and pale yellow for lower ones.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s4_1_2">
<label>4.1.2</label>
<title>Result of DenseNet121 model</title>
<p>
<xref ref-type="table" rid="T4">
<bold>Table&#xa0;4</bold>
</xref> illustrates that the retrained DenseNet121 model exhibits strong performance in both classes, as evidenced by the elevated accuracy, recall, and F1 Scores. For the Dubas class, the model achieves a precision of 99%, indicating that 97% of the examples predicted as Dubas are accurate, and a recall of 97%, demonstrating that the model recognizes 97% of all genuine Dubas occurrences. In the Healthy class, the accuracy and recall are 97% and 99%, respectively, indicating the model&#x2019;s proficiency in reliably classifying healthy samples. The F1-scores for both groups are 98%, indicating a balanced equilibrium between accuracy and recall. The model&#x2019;s overall accuracy of 98% highlights its success, accompanied by a weighted average precision, recall, and F1-score of 98%, demonstrating consistent performance across the dataset. The findings indicate that the DenseNet121 model has shown high performance.</p>
<table-wrap id="T4" position="float">
<label>Table&#xa0;4</label>
<caption>
<p>Result of DenseNet121 model.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Class name</th>
<th valign="middle" align="center">Precision (%)</th>
<th valign="middle" align="center">Recall (%)</th>
<th valign="middle" align="center">F1-score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">99</td>
<td valign="middle" align="center">97</td>
<td valign="middle" align="center">98</td>
</tr>
<tr>
<td valign="middle" align="center">Healthy</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">99</td>
<td valign="middle" align="center">98</td>
</tr>
<tr>
<td valign="middle" align="center">Accuracy</td>
<td valign="middle" align="center"/>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center"/>
</tr>
<tr>
<td valign="middle" align="center">Weighted_Avg</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">98</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>
<xref ref-type="fig" rid="f8">
<bold>Figure&#xa0;8</bold>
</xref> shows the confusion matrix for the DenseNet121 model in classifying healthy (1) and Dubas (0) data, demonstrating outstanding performance. Out of 160 test samples (80 in each class), the model correctly names 78 as Dubas and 79 as healthy. This led to three mistakes: two false positives (healthy samples were mistakenly labeled as Dubas) and one false negative (Dubas samples were labeled as Healthy). This yields an accuracy of 98.12%, demonstrating the efficacy of DenseNet121 in feature extraction and classification. The modestly reduced misclassification rate indicates a modest improvement in generalization.</p>
<fig id="f8" position="float">
<label>Figure&#xa0;8</label>
<caption>
<p>Confusion matrix DenseNet121model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g008.tif">
<alt-text content-type="machine-generated">Confusion matrix titled &#x201c;Test Set Confusion Matrix&#x201d; with actual values on the y-axis and predicted values on the x-axis. The matrix shows 78 true negatives, 79 true positives, 2 false positives, and 1 false negative. A color gradient from dark blue to light yellow represents data values.</alt-text>
</graphic>
</fig>
</sec>
<sec id="s4_1_3">
<label>4.1.3</label>
<title>Result of ViT model</title>
<p>
<xref ref-type="table" rid="T5">
<bold>Table&#xa0;5</bold>
</xref> presents the performance parameters of the ViT model in distinguishing between Dubas (0) and Healthy (1) samples, highlighting its robust predictive potential. The precision for Dubas is 100%, indicating that almost all samples categorized as Dubas are accurate. Meanwhile, the recall is 97%, demonstrating that 96% of genuine Dubas samples were accurately recognized. Correspondingly, for the Healthy class, the model achieves 98% accuracy and 100% recall, ensuring that most Healthy samples are accurately identified. The F1-score, which equilibrates accuracy and recall, is 99% for both classes, underscoring the model&#x2019;s overall dependability. The overall accuracy of 99% and a weighted average F1-score of 99% indicate that ViT exhibits uniform classification performance across both categories. The findings validate the model&#x2019;s efficacy in differentiating between Dubas and Healthy instances, exhibiting only negligible misclassifications, hence underscoring its robustness and generalizability in medical image classification.</p>
<table-wrap id="T5" position="float">
<label>Table&#xa0;5</label>
<caption>
<p>ViT model performance.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Class name</th>
<th valign="middle" align="center">Precision (%)</th>
<th valign="middle" align="center">Recall (%)</th>
<th valign="middle" align="center">F1-Score (%)</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">100</td>
<td valign="middle" align="center">97</td>
<td valign="middle" align="center">99</td>
</tr>
<tr>
<td valign="middle" align="center">Healthy</td>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center">100</td>
<td valign="middle" align="center">99</td>
</tr>
<tr>
<td valign="middle" align="center">Accuracy</td>
<td valign="middle" align="center"/>
<td valign="middle" align="center">98</td>
<td valign="middle" align="center"/>
</tr>
<tr>
<td valign="middle" align="center">Weighted_Avg_plam system</td>
<td valign="middle" align="center">99</td>
<td valign="middle" align="center">99</td>
<td valign="middle" align="center">99</td>
</tr>
</tbody>
</table>
</table-wrap>
<p>
<xref ref-type="fig" rid="f9">
<bold>Figure&#xa0;9</bold>
</xref> illustrates the confusion matrix for the ViT model in identifying Dubas (0) and Healthy (1) samples, indicating its excellent accuracy. The model accurately categorized 78 Dubas samples, misclassifying only 2 as Healthy, resulting in a high recall for the Dubas category. It accurately recognized all 80 Healthy samples without any false negatives, indicating that every true Healthy instance was found. This signifies that the model exhibits perfect recall for the Healthy class. However, its small misclassification in the Dubas group implies a little compromise in specificity. The ViT model demonstrates robust prediction performance with few mistakes, making it extremely dependable for differentiating between Dubas and Healthy instances.</p>
<fig id="f9" position="float">
<label>Figure&#xa0;9</label>
<caption>
<p>ViT model.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g009.tif">
<alt-text content-type="machine-generated">Confusion matrix titled &#x201c;Test Set Confusion Matrix&#x201d; showing predicted vs. actual values. True positives: 80, true negatives: 78. False positives: 0, false negatives: 2. Color gradient bar ranging from 0 to 80.</alt-text>
</graphic>
</fig>
</sec>
</sec>
<sec id="s4_2">
<label>4.2</label>
<title>Results performance</title>
<p>This segment of the research used cutting-edge deep learning and Vision Transformer models for the identification of palm leaf diseases. Previously trained on the ImageNet dataset, the publicly available Palm Leaves dataset was utilized to augment pre-trained deep learning (DL) and Vision Transformer (ViT) networks. Every model in our study was standardized using two output classes, a dropout rate of 0.5, and a learning rate of 0.001.</p>
<p>The dataset consisted of training, testing, and validation samples. Of the palm leaf samples, 80% were set aside for pre-training EfficientNetV2B0 models. Every model run for 10 epochs and showed that our model started to converge with high accuracy after this length. The first experiment demonstrates the performance of the EfficientNetV2B0 model, as illustrated in <xref ref-type="fig" rid="f10">
<bold>Figure&#xa0;10a</bold>
</xref>. <xref ref-type="fig" rid="f10">
<bold>Figure&#xa0;10b</bold>
</xref> displays the log loss of the EfficientNetV2B0 model. The EfficientNetV2B0 model reached a testing accuracy of 98.12%.</p>
<fig id="f10" position="float">
<label>Figure&#xa0;10</label>
<caption>
<p>Performance of the EfficientNetV2B model. <bold>(a)</bold> accuracy; <bold>(b)</bold> loss.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g010.tif">
<alt-text content-type="machine-generated">Graphs comparing model performance over epochs. Left graph shows accuracy, with training accuracy rising from 0.65 to about 0.90, and validation accuracy increasing slightly above 0.95. Right graph shows loss, with training and validation loss both decreasing from over 4.5 to around 1.5.</alt-text>
</graphic>
</fig>
<p>In the second experiment, we used the Palm dataset to test DenseNet-121. Based on <xref ref-type="fig" rid="f11">
<bold>Figure&#xa0;11a</bold>
</xref>, the model achieved a recognition accuracy of around 97.50% in the first 10 epochs, and then it increased to a high accuracy of 92.50%. At 0.20%, the recorded loss model and at 0.0974%, the validation model are shown in <xref ref-type="fig" rid="f11">
<bold>Figure&#xa0;11b</bold>
</xref>.</p>
<fig id="f11" position="float">
<label>Figure&#xa0;11</label>
<caption>
<p>Performance DenseNet121 model. <bold>(a)</bold> accuracy; <bold>(b)</bold> loss.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g011.tif">
<alt-text content-type="machine-generated">Two-line graphs depict model performance over epochs. Graph (a) shows accuracy, with training accuracy in blue and validation accuracy in green, both rising and stabilizing around 0.95. Graph (b) shows loss, with training loss in red and validation loss in yellow, both decreasing and stabilizing below 0.1.</alt-text>
</graphic>
</fig>
<p>Using the ViT model, the third experiment was carried out. The recognition accuracy graph and the validation and training loss graph are displayed in <xref ref-type="fig" rid="f12">
<bold>Figures&#xa0;12a, b</bold>
</xref>, respectively, and they demonstrate the identical methods used to evaluate model loss and recognition accuracy. The model&#x2019;s accuracy was 99.37%, with a margin of error of 0.01%.</p>
<fig id="f12" position="float">
<label>Figure&#xa0;12</label>
<caption>
<p>Performance ViT model. <bold>(a)</bold> accuracy; <bold>(b)</bold> loss.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g012.tif">
<alt-text content-type="machine-generated">Two graphs display model performance over 20 epochs. Graph (a), titled &#x201c;Accuracy,&#x201d; shows training accuracy in blue and validation accuracy in green, both consistently high. Graph (b), titled &#x201c;Loss,&#x201d; shows training loss in red and validation loss in yellow, both decreasing over time.</alt-text>
</graphic>
</fig>
<p>Optimal yields in agricultural production depend on the rapid diagnosis of crop diseases. The early detection of palm diseases using modern technologies is essential for maintaining an enhanced production rate. The literature review indicated that DL models excel in image classification, whereas DL methods effectively reduce training complexity and the need for extensive datasets. We examined three pre-trained models&#x2014;the EfficientNetV2-B0, the DenseNet-12, and the ViT models&#x2014;to determine which one was most effective in identifying various palm diseases. The pre-trained models were assessed using assessment criteria, including specificity, sensitivity, and F1 score values. <xref ref-type="table" rid="T6">
<bold>Table&#xa0;6</bold>
</xref> shows the results of the proposed systems against the existing system. We provided a visual representation of the validation accuracy for the pre-trained models by computing the validation accuracy using the F1 score. The ROC curve of the ViT model is shown in <xref ref-type="fig" rid="f13">
<bold>Figure&#xa0;13</bold>
</xref>. The model achieved a perfect score.</p>
<fig id="f13" position="float">
<label>Figure&#xa0;13</label>
<caption>
<p>ROC of ViTmodel.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g013.tif">
<alt-text content-type="machine-generated">ROC curve graph with a solid blue line indicating perfect classification with an AUC of 1.00. The dashed red line represents a random guess. The x-axis is the false positive rate and the y-axis is the true positive rate.</alt-text>
</graphic>
</fig>
<p>Using deep learning and ViT, <xref ref-type="fig" rid="f14">
<bold>Figure&#xa0;14</bold>
</xref> shows the method for plant leaf image categorization.</p>
<list list-type="simple">
<list-item>
<p>Step 1: Image Acquisition: A digital camera is used to capture plant leaves, both healthy and Unhealthy.</p>
</list-item>
<list-item>
<p>Step 2: Cloud Storage is used to centralize access to the images</p>
</list-item>
<list-item>
<p>Step 3: Pre-processing. Pre-processing of the system is used to handle resizing, normalization, and augmentation, thereby improving model training.</p>
</list-item>
<list-item>
<p>Step 4: Dataset Splitting - The system employs a validation process to divide the dataset into training, validation, and test sets, evaluating the model&#x2019;s performance.</p>
</list-item>
<list-item>
<p>Step 5: Model Training: The ViT and DL architectures are used to train the model on the training and validation datasets.</p>
</list-item>
<list-item>
<p>Step 6: Performance Evaluation - The trained model is evaluated on the test set, and its classification performance is visualized.</p>
</list-item>
<list-item>
<p>Step 7: Mobile Deployment: Farmers used the mobile applications for classifying plant leaves [<italic>Healthy</italic> vs. <italic>Dubas</italic>-infected] through a user-friendly mobile interface for field use.</p>
</list-item>
</list>
<fig id="f14" position="float">
<label>Figure&#xa0;14</label>
<caption>
<p>Deployment system for detecting Dubas insect diseases in palm leaves.</p>
</caption>
<graphic mimetype="image" mime-subtype="tiff" xlink:href="fpls-16-1612800-g014.tif">
<alt-text content-type="machine-generated">Flowchart illustrating a plant leaf image classification process. It starts with image acquisition, followed by cloud storage and image pre-processing. The dataset is split into test, training, and validation sets. Deep learning and Vision Transformer (ViT) models analyze the data. Results are classified as &#x201c;Healthy&#x201d; or &#x201c;Dubas&#x201d; and performance is assessed through a mobile application interface.</alt-text>
</graphic>
</fig>
<table-wrap id="T6" position="float">
<label>Table&#xa0;6</label>
<caption>
<p>Comparative performance analysis of various network models.</p>
</caption>
<table frame="hsides">
<thead>
<tr>
<th valign="middle" align="center">Authors, year</th>
<th valign="middle" align="center">Infested palm</th>
<th valign="middle" align="center">Model</th>
<th valign="middle" align="center">Acc %</th>
</tr>
</thead>
<tbody>
<tr>
<td valign="middle" align="center">
<xref ref-type="bibr" rid="B3">Al-Mulla et&#xa0;al., 2023</xref>
</td>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">CNN</td>
<td valign="middle" align="center">93-95%</td>
</tr>
<tr>
<td valign="middle" align="center">
<xref ref-type="bibr" rid="B4">Alshehhi et&#xa0;al., 2022</xref>
</td>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">CNN, GoogleNet</td>
<td valign="middle" align="center">98%</td>
</tr>
<tr>
<td valign="middle" align="center">Proposed ViT</td>
<td valign="middle" align="center">Dubas</td>
<td valign="middle" align="center">ViT</td>
<td valign="middle" align="left">99.37%</td>
</tr>
</tbody>
</table>
</table-wrap>
</sec>
</sec>
<sec id="s5" sec-type="conclusions">
<label>5</label>
<title>Conclusion</title>
<p>This study investigated a transfer learning and transformer methodology to provide an edge computing solution for the identification and detection of palm leaf diseases. Python supports three pre-trained models. This study presents a novel method for automatically detecting palm leaf disease against a natural background. This enables the differentiation between groups of infected and healthy leaves. We developed the technique for EfficientNetV2B0, DenseNet12, and the transformer paradigm. The dataset comprises 1600 images of palm leaves, with 800 depicting healthy and 800 depicting Dubas. To reduce computing time, pre-processing was performed, including image resizing and normalization, followed by augmentation. Augmentation was implemented by rotation, flipping, shearing, and zooming methods. The models were used to identify palm leaf disease using the TensorFlow framework with an input dimension of 224 &#xd7; 224 &#xd7; 3. The suggested approach achieved superior performance, as indicated by an accuracy value. The experiment demonstrated that the ViT model outperformed the other three models, achieving a validation accuracy of 99.37%, which is comparable to previously published techniques. The developed model successfully sustained elevated recall values, accuracy, and F1 scores. Although several automated detection models for palm leaf disease have been developed, their efficacy has often proven insufficient due to the resemblance of class attributes. This study primarily focused on detecting Dubas insect-related diseases and healthy leaves. This limitation of the study did not include other diseases, such as Brittle Leaves and Brown Leaf Spot.</p>
</sec>
</body>
<back>
<sec id="s6" sec-type="data-availability">
<title>Data availability statement</title>
<p>The datasets presented in this study can be found in online repositories. The names of the repository/repositories and accession number(s) can be found below: <uri xlink:href="https://www.kaggle.com/datasets/warcoder/palm-leaves-dataset">https://www.kaggle.com/datasets/warcoder/palm-leaves-dataset</uri>.</p>
</sec>
<sec id="s7" sec-type="author-contributions">
<title>Author contributions</title>
<p>TA: Writing &#x2013; review &amp; editing, Writing &#x2013; original draft, Methodology, Visualization, Formal Analysis, Conceptualization, Data curation, Project administration, Investigation. HA: Writing &#x2013; review &amp; editing, Software, Validation, Funding acquisition, Resources, Formal Analysis, Writing &#x2013; original draft, Methodology, Supervision, Investigation, Project administration.</p>
</sec>
<sec id="s8" sec-type="funding-information">
<title>Funding</title>
<p>The author(s) declare financial support was received for the research and/or publication of this article. This study has been funded by Date Palm Research Center of Excellence, King Faisal University, Saudi Arabia, through funding the research project number DPRC-11-2024.</p>
</sec>
<sec id="s9" sec-type="COI-statement">
<title>Conflict of interest</title>
<p>The authors declare that the research was conducted in the absence of any commercial or financial relationships that could be construed as a potential conflict of interest.</p>
</sec>
<sec id="s10" sec-type="ai-statement">
<title>Generative AI statement</title>
<p>The author(s) declare that no Generative AI was used in the creation of this manuscript.</p>
<p>Any alternative text (alt text) provided alongside figures in this article has been generated by Frontiers with the support of artificial intelligence and reasonable efforts have been made to ensure accuracy, including review by the authors wherever possible. If you identify any issues, please contact us.</p>
</sec>
<sec id="s11" sec-type="disclaimer">
<title>Publisher&#x2019;s note</title>
<p>All claims expressed in this article are solely those of the authors and do not necessarily represent those of their affiliated organizations, or those of the publisher, the editors and the reviewers. Any product that may be evaluated in this article, or claim that may be made by its manufacturer, is not guaranteed or endorsed by the publisher.</p>
</sec>
<ref-list>
<title>References</title>
<ref id="B1">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Abu-zanona</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Elaiwat</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Younis</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Innab</surname> <given-names>N.</given-names>
</name>
<name>
<surname>Kamruzzaman</surname> <given-names>M. M.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Classification of palm trees diseases using convolution neural network</article-title>. <source>Int. J. Adv. Comput. Sci. Appl.</source> <volume>13</volume>, <fpage>943</fpage>&#x2013;<lpage>949</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.14569/IJACSA.2022.01306111</pub-id>
</citation></ref>
<ref id="B2">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Almadhor</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Rauf</surname> <given-names>H. T.</given-names>
</name>
<name>
<surname>Lali</surname> <given-names>M. I.</given-names>
</name>
<name>
<surname>Dama&#x161;evi&#x10d;ius</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Alouffi</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Alharbi</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>AI-driven framework for recognition of guava plant diseases through machine learning from DSLR camera sensor based high resolution imagery</article-title>. <source>Sensors</source> <volume>21</volume>, <elocation-id>3830</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/s21113830</pub-id>, PMID: <pub-id pub-id-type="pmid">34205885</pub-id></citation></ref>
<ref id="B3">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Al-Mulla</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Ali</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Parimi</surname> <given-names>K.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Detection and analysis of dubas-infested date palm trees using deep learning, remote sensing, and GIS techniques in wadi bani kharus</article-title>. <source>Sustainability</source> <volume>15</volume>, <elocation-id>14045</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/su151914045</pub-id>
</citation></ref>
<ref id="B4">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Alshehhi</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Almannaee</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Shatnawi</surname> <given-names>M.</given-names>
</name>
</person-group> (<year>2022</year>). &#x201c;<article-title>Date palm leaves discoloration detection system using deep transfer learning</article-title>,&#x201d; in <source>Proc. Int. Conf. Emerg. Technol. Intell. Syst</source> (<publisher-name>Springer</publisher-name>, <publisher-loc>Cham, Switzerland</publisher-loc>), <fpage>150</fpage>&#x2013;<lpage>161</lpage>.</citation></ref>
<ref id="B5">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sharma</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Jain</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Gupta</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Chowdary</surname> <given-names>V.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Machine learning applications for precision agriculture: A comprehensive review</article-title>. <source>IEEE Access</source> <volume>9</volume>, <fpage>4843</fpage>&#x2013;<lpage>4873</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/Access.6287639</pub-id>
</citation></ref>
<ref id="B6">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ashqar</surname> <given-names>B. A. M.</given-names>
</name>
<name>
<surname>Abu-Naser</surname> <given-names>S. S.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Image-based tomato leaves diseases detection using deep learning</article-title>. <source>Int. J. Acad. Eng. Res.</source> <volume>2</volume>, <fpage>10</fpage>&#x2013;<lpage>16</lpage>.</citation></ref>
<ref id="B7">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Atila</surname> <given-names>&#xdc;.</given-names>
</name>
<name>
<surname>U&#xe7;ar</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Akyol</surname> <given-names>K.</given-names>
</name>
<name>
<surname>U&#xe7;ar</surname> <given-names>E.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Plant leaf disease classification using EfficientNet deep learning model</article-title>. <source>Ecol. Inf.</source> <volume>61</volume>, <elocation-id>101182</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.ecoinf.2020.101182</pub-id>
</citation></ref>
<ref id="B8">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ba-Angood</surname> <given-names>S. A.</given-names>
</name>
<name>
<surname>Al-Ghurabi</surname> <given-names>A. S.</given-names>
</name>
<name>
<surname>Hubaishan</surname> <given-names>M. A.</given-names>
</name>
</person-group> (<year>2009</year>). <article-title>Biology and chemical control of the old world bug (Doubas bug) Ommatissus lybicus DeBerg on date palm trees in the coastal areas of Hadramout Governorate, Republic of Yemen</article-title>. <source>Arab J. Plant Prot.</source> <volume>27</volume>, <fpage>1</fpage>&#x2013;<lpage>9</lpage>.</citation></ref>
<ref id="B9">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Bayomi-Alli</surname> <given-names>O. O.</given-names>
</name>
<name>
<surname>Dama&#x161;evi&#x10d;ius</surname> <given-names>R.</given-names>
</name>
<name>
<surname>Misra</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Maskeli&#x16b;nas</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Cassava disease recognition from low-quality images using enhanced data augmentation model and deep learning</article-title>. <source>Expert Syst.</source> <volume>38</volume>, <fpage>1</fpage>&#x2013;<lpage>12</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/exsy.12746</pub-id>
</citation></ref>
<ref id="B10">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Elwan</surname> <given-names>A. A.</given-names>
</name>
<name>
<surname>Al-Tamimi</surname> <given-names>S. S.</given-names>
</name>
</person-group> (<year>1999</year>). <article-title>Life cycle of Dubas bug Ommatissus binotatus lybicus de Berg. (Homoptera: Tropiduchidae) in Sultanate of Oman</article-title>. <source>Egypt. J. Agric. Res.</source> <volume>77</volume>, <fpage>1547</fpage>&#x2013;<lpage>1553</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.21608/ejar.1999.342384</pub-id>
</citation></ref>
<ref id="B11">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Eunice</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Popescu</surname> <given-names>D. E.</given-names>
</name>
<name>
<surname>Chowdary</surname> <given-names>M. K.</given-names>
</name>
<name>
<surname>Hemanth</surname> <given-names>J.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Deep learning-based leaf disease detection in crops using images for agricultural applications</article-title>. <source>Agronomy</source> <volume>12</volume>, <fpage>2395</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/dme.15222</pub-id>, PMID: <pub-id pub-id-type="pmid">37690127</pub-id></citation></ref>
<ref id="B12">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hamdani</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Septiarini</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Sunyoto</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Suyanto</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Utaminingrum</surname> <given-names>F.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Detection of oil palm leaf disease based on color histogram and supervised classifier</article-title>. <source>Optik (Stuttg)</source> <volume>245</volume>, <elocation-id>167753</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.ijleo.2021.167753</pub-id>
</citation></ref>
<ref id="B13">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Harakannanavar</surname> <given-names>S. S.</given-names>
</name>
<name>
<surname>Rudagi</surname> <given-names>J. M.</given-names>
</name>
<name>
<surname>Puranikmath</surname> <given-names>V. I.</given-names>
</name>
<name>
<surname>Siddiqua</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Pramodhini</surname> <given-names>R.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Plant leaf disease detection using computer vision and machine learning algorithms</article-title>. <source>Global Transition Proc.</source> <volume>3</volume>, <fpage>305</fpage>&#x2013;<lpage>310</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.gltp.2022.03.016</pub-id>
</citation></ref>
<ref id="B14">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Hessane</surname> <given-names>A.</given-names>
</name>
<name>
<surname>El Youssefi</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Farhaoui</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Aghoutane and F. Amounas</surname> <given-names>B.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>A machine learning based framework for a stage-wise classification of date palm white scale disease</article-title>. <source>Big Data Min. Anal.</source> <volume>6</volume>, <fpage>263</fpage>&#x2013;<lpage>272</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.26599/BDMA.2022.9020022</pub-id>
</citation></ref>
<ref id="B15">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Iqbal</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Hussain</surname> <given-names>I.</given-names>
</name>
<name>
<surname>Hakim</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Ullah</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Yousuf</surname> <given-names>H. M.</given-names>
</name>
</person-group> (<year>2023</year>). <article-title>Early detection and classification of rice brown spot and bacterial blight diseases using digital image processing</article-title>. <source>J. Computing Biomed. Inf.</source> <volume>4</volume>, <fpage>98</fpage>&#x2013;<lpage>109</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1111/dme.15222</pub-id>, PMID: <pub-id pub-id-type="pmid">37690127</pub-id></citation></ref>
<ref id="B16">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Jaradat</surname> <given-names>A. A.</given-names>
</name>
</person-group> (<year>2015</year>). &#x201c;<article-title>Biodiversity, genetic diversity, and genetic resources of date palm</article-title>,&#x201d; in <source>Date Palm Genetic Resources and Utilization: Volume 1: Africa and the Americas</source>. Eds. <person-group person-group-type="editor">
<name>
<surname>Al-Khayri</surname> <given-names>J. M.</given-names>
</name>
<name>
<surname>Jain</surname> <given-names>S. M.</given-names>
</name>
<name>
<surname>Johnson</surname> <given-names>D. V.</given-names>
</name>
</person-group> (<publisher-name>Springer</publisher-name>, <publisher-loc>Dordrecht, The Netherlands</publisher-loc>), <fpage>19</fpage>&#x2013;<lpage>71</lpage>.</citation></ref>
<ref id="B17">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kamal</surname> <given-names>M. M.</given-names>
</name>
<name>
<surname>Masazhar</surname> <given-names>A. N. I.</given-names>
</name>
<name>
<surname>Rahman</surname> <given-names>F. A.</given-names>
</name>
</person-group> (<year>2018</year>). <article-title>Classification of leaf disease from image processing technique</article-title>. <source>Indones. J. Electr.Eng. Comput. Sci.</source> <volume>10</volume>, <fpage>191</fpage>&#x2013;<lpage>200</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.11591/ijeecs.v10.i1.pp191-200</pub-id>
</citation></ref>
<ref id="B18">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Khamparia</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Saini</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Gupta</surname> <given-names>D.</given-names>
</name>
<name>
<surname>Khanna</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Tiwari</surname> <given-names>S.</given-names>
</name>
<name>
<surname>de Albuquerque</surname> <given-names>V. H. C.</given-names>
</name>
<etal/>
</person-group>. (<year>2020</year>). <article-title>Seasonal crops disease prediction and classifcation using deep convolutional encoder network</article-title>. <source>Circuits Syst. Signal Process</source> <volume>39</volume>, <fpage>818</fpage>&#x2013;<lpage>836</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s00034-019-01041-0</pub-id>
</citation></ref>
<ref id="B19">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Kong</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Wang</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Yang</surname> <given-names>C.</given-names>
</name>
<name>
<surname>Jin</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Zuo</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Zhang</surname> <given-names>X.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>A spatial feature-enhanced attention neural network with high-order pooling representation for application in pest and disease recognition</article-title>. <source>Agriculture</source> <volume>12</volume>, <fpage>500</fpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/agriculture12040500</pub-id>
</citation></ref>
<ref id="B20">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Li</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Yan</surname> <given-names>P.</given-names>
</name>
<name>
<surname>Li</surname> <given-names>G.</given-names>
</name>
<name>
<surname>Jiang</surname> <given-names>Y.</given-names>
</name>
<name>
<surname>Luo</surname> <given-names>H.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Deep learning attention mechanism in medical image analysis: Basics and beyonds</article-title>. <source>Int. J. Netw. Dyn Intell.</source> <volume>2023</volume>, <fpage>93</fpage>&#x2013;<lpage>116</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.53941/ijndi0201006</pub-id>
</citation></ref>
<ref id="B21">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Liu</surname> <given-names>X.</given-names>
</name>
<name>
<surname>Ghazali</surname> <given-names>K. H.</given-names>
</name>
<name>
<surname>Han</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Mohamed</surname> <given-names>I. I.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Automatic detection of oil palm tree from UAV images based on the deep learning method</article-title>. <source>Appl. Artif. Intell.</source> <volume>35</volume>, <fpage>13</fpage>&#x2013;<lpage>24</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1080/08839514.2020.1831226</pub-id>
</citation></ref>
<ref id="B22">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Masazhar</surname> <given-names>A. N. I.</given-names>
</name>
<name>
<surname>Kamal</surname> <given-names>M. M.</given-names>
</name>
</person-group> (<year>2017</year>). &#x201c;<article-title>Digital image processing technique for palm oil leaf disease detection using multiclass SVM classifier</article-title>,&#x201d; <source>2017 IEEE 4th International Conference on Smart Instrumentation, Measurement and Application (ICSIMA)</source>, (<publisher-loc>Putrajaya, Malaysia</publisher-loc>) pp. <fpage>1</fpage>&#x2013;<lpage>6</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICSIMA.2017.8311978</pub-id>
</citation></ref>
<ref id="B23">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Nusrat</surname> <given-names>J.</given-names>
</name>
<name>
<surname>Paulo</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Zhaohui</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Andrew</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Jithin</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Zhao</surname> <given-names>Z.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Detecting and distinguishing wheat diseases using image processing and machine learning Algorithms</article-title>. <source>2020 ASABE Annu. Int. Virtual Meeting</source>. doi:&#xa0;<pub-id pub-id-type="doi">10.13031/aim.20200037</pub-id>
</citation></ref>
<ref id="B24">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Payandeh</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Kamali</surname> <given-names>K.</given-names>
</name>
<name>
<surname>Fathipour</surname> <given-names>Y.</given-names>
</name>
</person-group> (<year>2010</year>). <article-title>Population structure and seasonal activity of Ommatissus lybicus in Bam region of Iran (Homoptera Tropiduchidae)</article-title>. <source>Munis Entomol. Zool.</source> <volume>5</volume>, <fpage>726</fpage>&#x2013;<lpage>733</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1038/nature10238</pub-id>, PMID: <pub-id pub-id-type="pmid">21743477</pub-id></citation></ref>
<ref id="B25">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Resh</surname> <given-names>V. H.</given-names>
</name>
<name>
<surname>Card&#xe9;</surname> <given-names>R. T.</given-names>
</name>
</person-group> (<year>2009</year>). <source>Encyclopedia of Insects</source>. <edition>2nd</edition> (<publisher-loc>London, UK</publisher-loc>: <publisher-name>Academic Press</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>1136</lpage>.</citation></ref>
<ref id="B26">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Septiarini</surname> <given-names>A.</given-names>
</name>
<name>
<surname>Hamdani</surname> <given-names>H.</given-names>
</name>
<name>
<surname>Hardianti</surname> <given-names>T.</given-names>
</name>
<name>
<surname>Winarno</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Suyanto and E. Irwansyah</surname> <given-names>S.</given-names>
</name>
</person-group> (<year>2021</year>). <article-title>Pixel quantification and color feature extraction on leaf images for oil palm disease identification</article-title>. <source>Proc. 7th Int. Conf. Electr. Electron. Inf. Eng. (ICEEIE)</source>, <volume>39</volume>, <fpage>1</fpage>&#x2013;<lpage>5</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICEEIE52663.2021.9616645</pub-id>
</citation></ref>
<ref id="B27">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Sethy</surname> <given-names>P. K.</given-names>
</name>
<name>
<surname>Barpanda</surname> <given-names>N. K.</given-names>
</name>
<name>
<surname>Rath</surname> <given-names>A. K.</given-names>
</name>
<name>
<surname>Behera</surname> <given-names>S. K.</given-names>
</name>
</person-group> (<year>2020</year>). <article-title>Image processing techniques for diagnosing rice plant disease: a survey</article-title>. <source>Proc. Comput. Sci.</source> <volume>167</volume>, <fpage>516</fpage>&#x2013;<lpage>530</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.procs.2020.03.308</pub-id>
</citation></ref>
<ref id="B28">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Shah</surname> <given-names>J. P.</given-names>
</name>
<name>
<surname>Prajapati</surname> <given-names>H. B.</given-names>
</name>
<name>
<surname>Dabhi</surname> <given-names>V. K.</given-names>
</name>
</person-group> (<year>2016</year>). &#x201c;<article-title>A survey on detection and classification of rice plant diseases</article-title>,&#x201d; in <source>2016 IEEE International Conference on Current Trends in Advanced Computing (ICCTAC)</source> (<publisher-loc>Bangalore, India</publisher-loc>: <publisher-name>IEEE</publisher-name>), <fpage>1</fpage>&#x2013;<lpage>8</lpage>.</citation></ref>
<ref id="B29">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Shruthi</surname> <given-names>U.</given-names>
</name>
<name>
<surname>Nagaveni</surname> <given-names>V.</given-names>
</name>
<name>
<surname>Raghavendra</surname> <given-names>B. K.</given-names>
</name>
</person-group> (<year>2019</year>). &#x201c;<article-title>A Review on Machine Learning Classification Techniques for Plant Disease Detection</article-title>,&#x201d; <source>2019 5th International Conference on Advanced Computing &amp; Communication Systems (ICACCS)</source>, (<publisher-loc>Coimbatore, India</publisher-loc>) <fpage>281</fpage>&#x2013;<lpage>284</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1109/ICACCS.2019.8728415</pub-id>
</citation></ref>
<ref id="B30">
<citation citation-type="confproc">
<person-group person-group-type="author">
<name>
<surname>Thacker</surname> <given-names>J. R. M.</given-names>
</name>
<name>
<surname>Al-Mahmooli</surname> <given-names>I. H. S.</given-names>
</name>
<name>
<surname>Deadman</surname> <given-names>M. L.</given-names>
</name>
</person-group> (<year>2003</year>). &#x201c;<article-title>Population dynamics and control of the dubas bug Ommatissus lybicus in the Sultanate of Oman. In Proceeding of the BCPC International Congress: Crop Science and Technology, Volumes 1 and 2</article-title>,&#x201d; in <conf-name>Proceedings of an International Congress Held at the SECC</conf-name>, (<publisher-loc>Alton, Hampshire, UK</publisher-loc>: <publisher-name>British Crop Protection Council (BCPC)</publisher-name>) <conf-date>10&#x2013;12 November</conf-date>.</citation></ref>
<ref id="B31">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>The Ministry of Environment</collab>
<collab>Water and Agriculture Saudi Arabia</collab>
</person-group> (<year>2020</year>).<article-title>FAO approves Saudi Arabia&#x2019;s proposal to declare 2027 the international year of date palm</article-title>. Available online at: <uri xlink:href="https://mewa.gov.sa/en/MediaCenter/News/Pages/News201220.aspx/">https://mewa.gov.sa/en/MediaCenter/News/Pages/News201220.aspx/</uri> (Accessed <access-date>April 01, 2024</access-date>).</citation></ref>
<ref id="B32">
<citation citation-type="web">
<person-group person-group-type="author">
<collab>The Saudi Press Agency (SPA)</collab>
</person-group> (<year>2024</year>).<article-title>26,000 date farms in Madinah yield SAR948.5 million in 2023</article-title>. Available online at: <uri xlink:href="https://www.spa.gov.sa/en/N2035544">https://www.spa.gov.sa/en/N2035544</uri> (Accessed <access-date>April 01, 2024</access-date>).</citation></ref>
<ref id="B33">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Tholkapiyan</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Aruna Devi</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Bhatt</surname> <given-names>B.</given-names>
</name>
<name>
<surname>Saravana Kumar</surname> <given-names>E.</given-names>
</name>
<name>
<surname>Kirubakaran</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Kumar</surname> <given-names>R.</given-names>
</name>
<etal/>
</person-group>. (<year>2023</year>). <article-title>Performance analysis of rice plant diseases identification and classification methodology</article-title>. <source>Wireless Pers. Commun.</source> <volume>130</volume>, <fpage>1317</fpage>&#x2013;<lpage>1341</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1007/s11277-023-10333-3</pub-id>
</citation></ref>
<ref id="B34">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Too</surname> <given-names>E. C.</given-names>
</name>
<name>
<surname>Yujian</surname> <given-names>L.</given-names>
</name>
<name>
<surname>Njuki</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Yingchun</surname> <given-names>L.</given-names>
</name>
</person-group> (<year>2019</year>). <article-title>A comparative study of fine-tuning deep learning models for plant disease identification</article-title>. <source>Comput. Electron. Agric.</source> <volume>161</volume>, <fpage>272</fpage>&#x2013;<lpage>279</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.1016/j.compag.2018.03.032</pub-id>
</citation></ref>
<ref id="B35">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Ya&#x11f;</surname> <given-names>&#x130;.</given-names>
</name>
<name>
<surname>Altan</surname> <given-names>A.</given-names>
</name>
</person-group> (<year>2022</year>). <article-title>Artificial intelligence-based robust hybrid algorithm design and implementation for real-time detection of plant diseases in agricultural environments</article-title>. <source>Biology</source> <volume>11</volume>, <elocation-id>1732</elocation-id>. doi:&#xa0;<pub-id pub-id-type="doi">10.3390/biology11121732</pub-id>, PMID: <pub-id pub-id-type="pmid">36552243</pub-id></citation></ref>
<ref id="B36">
<citation citation-type="book">
<person-group person-group-type="author">
<name>
<surname>Zaid</surname> <given-names>A.</given-names>
</name>
<name>
<surname>de Wet</surname> <given-names>P.</given-names>
</name>
</person-group> (<year>2002</year>). &#x201c;<article-title>Origin, geographical distribution and nutritional values of date palm</article-title>,&#x201d; in <source>Date Palm Cultivation</source>. Ed. <person-group person-group-type="editor">
<name>
<surname>Zaid</surname> <given-names>A.</given-names>
</name>
</person-group> (<publisher-name>Food and Agriculture Organization of the United Nations (FAO</publisher-name>, <publisher-loc>Rome, Italy</publisher-loc>). Chapter II.</citation></ref>
<ref id="B37">
<citation citation-type="journal">
<person-group person-group-type="author">
<name>
<surname>Zia Ur Rehman</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Ahmed</surname> <given-names>F.</given-names>
</name>
<name>
<surname>Attique Khan</surname> <given-names>M.</given-names>
</name>
<name>
<surname>Tariq</surname> <given-names>U.</given-names>
</name>
<name>
<surname>Shaukat Jamal</surname> <given-names>S.</given-names>
</name>
<name>
<surname>Ahmad</surname> <given-names>J.</given-names>
</name>
<etal/>
</person-group>. (<year>2022</year>). <article-title>Classification of citrus plant diseases using deep transfer learning</article-title>. <source>Comput. Mater. Continua</source> <volume>70</volume>, <fpage>1401</fpage>&#x2013;<lpage>1417</lpage>. doi:&#xa0;<pub-id pub-id-type="doi">10.32604/cmc.2022.019046</pub-id>
</citation></ref>
</ref-list>
</back>
</article>