@article{ 
author = {Maskanati, Salman and Keshavarz, Ahm},  
title = {Online Persian Hand Writing Recognition Using Language Model and Reduction of User Writing Rules}, 
abstract ={The Joint-up, cursive form of Persian words and immense variety of its scripts, also different figures of Persian letters depending on their sitting positions in the words, have turned the Persian handwritings recognition to an intense challenge. The major obstacle of the most often recognition ways, is their inattention to sentence contexture which causes utilizing of a word with correct appearance within an incorrect sentence, when an input word is misrecognized. Sketching a solution that provides suitable analysis of sentence contexture, requires huge linguistic resources to take place as a fine representative for the chosen language to be recognized. In this article, a new method for online recognition of Persian words is presented which tries to improve recognition process by using the term contexture. In this article, the vocabularies collection of Persian language is divided into two groups. The first category is the vocabulary with all of their sub-words being supported by the database of handwritten subclasses, while these vocabulary form 68.2% of the total vocabulary, and the assumptions being scored at the recognition stage, are members of these vocabularies. The second category is the vocabulary that is not supported by the database. Obviously, if the recognition system does not support this vocabulary, it cannot recognize more than 30 percentages of the language&#39;s words. At the recognition stage, the symptoms are detected and a symptom tag is produced. Also, at this stage, using the same label, the vocabulary is also selected as the sign with the input word. (These vocabularies are chosen from those were not supported at the recognition stage). Scoring for hypotheses was done by combining recognition scores and linguistic models. The certain fact in this section is that it is impossible to calculate recognition scores due to the absence of hypothetical subheadings. Therefore, the vocabulary score being recognized in the previous steps, is used. According to the studies, it was concluded that if the word is equivalent to a member&#39;s input from a supported vocabulary, even if the result of the recognition is incorrect, in most cases the correct term is in the first four hypotheses. Usually, scores of the first few hypotheses are close to each other, and the other assumptions are far from the correct hypothesis. Since the system operates online, unnecessary computations should be avoided. Therefore, if the number of hypotheses in the recognition section are more than four hypotheses, only the first four hypotheses are calculated for the language model. To calculate the recognition score for new hypotheses, if there are fewer than four hypotheses in the recognition section, the lowest hypothesis score and otherwise the hypothesis score are considered for the recognition score of the new hypotheses. Then, as with previous assumptions, for the new hypotheses, the linguistic score is calculated, and then the final score is obtained for each hypothesis. Finally, the assumption with the highest score is considered as the system output, and the rest of the assumptions are displayed in the output to the user. Experiments show that even in the event of a mistake, the correct word is often presented as a second hypothesis in most cases, and in some cases as a third hypothesis. Also, to reduce the limits and rules that gainers compel to submit. The method demonstrated in this article includes the symptoms and morphemes framework of input handwritten are segregated and the framework of each morpheme with its symptoms is specified at first, then the symptoms of morphemes are specified and based on them a collection of words is being considered as a hypothesis. Each hypothesis is given a score by measuring the similarity to input handwritten and according to taken scores, the likely hypotheses are indicated. Then, this procedure is led to achieve hypotheses more likely by lingual models. To totalize the scores of a hypothesis, for the differences in scale of taken scores, a method of score normalization is being offered. The results demonstrate that by utilizing of a language model with an online system of handwriting recognition, a significant reduction of words recognition error rate is being achieved. In addition to error rate reduction, by taking advantages of this language model, a technique is being offered that can handle the Persian vocabulary recognition entirely. By availing the offered manner, the recognition precision at initial stage of letters level up to 95.9% and so the language model recognition up to 99.3% improved. So, using huge linguistic resources for Persian language and utilizing a language model, can improve the accuracy of recognition. For further work, reinforcement learning algorithm is suggested to adapt the algorithm for users. &#160;},  
Keywords = {Online Recognition, Persian Handwriting, k-nearest Neighbor, Language Model, User Limitation },
volume = {14},
Number = {2}, 
pages = {3-24}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {تشخیص دست‌نوشتۀ‌ برخط فارسی با استفاده از مدل زبانی و کاهش قوانین نگارش کاربر},
abstract_fa ={پیوسته&#8204;بودن کلمات فارسی و وجود تنوع بسیار زیاد رسم&#8204;الخط این زبان و همچنین شکل&#8204;های متنوع حروف فارسی بسته به محل قرارگیری&#8204;شان در کلمه، تشخیص دست&#8204;نوشته&#8204;های فارسی را به چالش کشانده&#8204;اند. مهم&#8204;ترین اشکال در اغلب روش&#8204;های بازشناسی بی&#8204;توجهی به بافت جمله است که باعث می&#8204;شود در مواردی که کلمه ورودی اشتباه بازشناسی می&#8204;شود، واژه&#8204;ای با ظاهر درست در جمله&#8204;ای نابه&#8204;جا به کار رود. طراحی مدلی که بتواند بافت جمله را به&#8204;خوبی تحلیل کند، مستلزم در&#8204;اختیار&#8204;داشتن منابع زبانی حجیمی است که نمایندۀ خوبی از زبان مورد بازشناسی باشند. در این مقاله روش جدیدی برای بازشناسی کلمات برخط فارسی ارائه شده است که با استفاده از بافت جمله سعی در بهبود بازشناسی دارد. فرآیند بازشناسی معرفی&#8204;شده در این نوشتار به این صورت است که ابتدا علائم و بدنه زیرکلمات دست&#8204;نوشته ورودی تفکیک شده و بدنه هر زیرکلمه و علائم آن مشخص می&#8204;شود؛ سپس علائم زیرکلمات تشخیص داده&#8204;شده و بر اساس آن مجموعه&#8204;ای از واژگان به&#8204;عنوان فرضیه در نظر گرفته می&#8204;شوند؛ به هر فرضیه بر اساس میزان شباهت آن به دست&#8204;نوشته ورودی امتیازی تعلق می&#8204;گیرد و بر اساس امتیاز حاصله محتمل&#8204;ترین فرضیات مشخص می&#8204;شوند. سپس این رویه توسط مدل زبانی برای یافتن فرضیات محتمل&#8204;تر، هدایت می&#8204;شود. نتایج آزمایش&#8204;های به&#8204;عمل&#8204;آمده نشان می&#8204;دهد که کاهش قابل توجهی در نرخ خطای بازشناسی کلمات حاصل شده و کاربر در نگارش ملزم به رعایت محدودیت&#8204;های کمتری است. از طرفی روش پیشنهادی می&#8204;تواند نسبت به روش&#8204;های قبلی با در&#8204;اختیار&#8204;داشتن یک پایگاه داده دست&#8204;نویس محدود، صحت مطلوب&#8204;تری ارائه کند. با به&#8204;کارگیری روش ارائه&#8204;شده، دقت بازشناسی در مرحلۀ&#8204; اولیه در سطح حروف 9/95% و پس از بازشناسی به&#8204;کمک مدل زبانی دقت بازشناسی به 3/99% ارتقا یافت. برای بهبود عملکرد الگوریتم، استفاده از الگوریتم یادگیری تقویتی برای تطبیق پذیری الگوریتم با نویسنده به&#8204;عنوان کار آینده پیشنهاد می&#8204;شود. &#160;},
keywords_fa = {بازشناسی برخط, دست‌نوشته فارسی, نزدیک‌ترین همسایه, مدل زبانی, محدودیت کاربر},

doi = {10.18869/acadpub.jsdp.14.2.3},
url = {http://jsdp.rcisp.ac.ir/article-1-428-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-428-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Ghasemi, Jamal and Kord, Somayeh and Gholami, Moham},  
title = {Classification of Cardiac Arrhythmias based on combination of the results of Neural Networks using Dempster-Shefer Evidence Theory}, 
abstract ={Cardiac arrhythmias are one of the most common heart diseases that may cause the death of the patient. Therefore, it is extremely important to detect cardiac arrhythmias.&#160; 3 categories of arrhythmia, namely, PAC, PVC, and normal are considered in this paper based on classifier fusion using evidence theory. In this study, at first a sample is carrying out the ECG signal with 250 point.&#160; Moreover, in each of the sampling, the maximum values will be obtained. Then, the average of the calculated values would be considered as adaptive thresholding and the total signals are multiplied by the inverse adaptive thresholding. After fixing the adaptive thresholding at number one, total resulting signal is becoming the power of 2. In this situation, the amounts smaller than one, are weakened and the larger than one amounts are reinforced. The smaller amount is removed and other amounts are held. Then, the maximum in each of the sampling is considered. In sampling areas that there is no peak, some maximum can be identified with zero value that these points should be removed from the set maximum. To find the maximum point where the maximum is close to the borders of sampling, two peaks may be placed in one field. This problem leads to removing one peak and non-recognition of the smaller peak. Some peaks near the border of sampling, for example the previous or next point on the border may be identified as the peak which eliminates the major peak and identifies the unrealistic peak. To solve this problem, the 80-point sampling is performed around each detected peak and the maximum value is obtained at the sampling areas. In this way, the correct peaks are identified and the wrong one will be deleted. In some parts, the peak signal is not quite sharp, and maybe two or more points that are adjacent to each other with the same value, will be considered as a peak. In other words, a closed peak is detected several times, which leads to detection of extra and incorrect peaks. In these circumstances, according to an amount that only belongs to one peak, just one of them should be considered and the other should be removed. After these steps, an obtained signal which includes peaks R, is compared with the original signal. To achieve the correct answer, it changes the number of sampling points and each time the result is compared with the previous values and with the original signal, too, until finally the major peaks will be identified. Then, HRV signal be will calculated. Linear properties contain root mean square of successive differences between normal intervals (RMSSD) and standard deviation of normal to normal intervals in a row (SDNN) and also heart rate (HR Mean) are calculated. Around each peak, 81 points window is inserted. These points for each peak is in one row. So resulting matrix (X) has 81 columns and its rows are the number of R peaks. SVD of matrix(X) is calculated. The obtained Matrix S will include the individual values. These singular signal values are non-linear features. If all used values are single, they can eclipse the linear features which will lead to the lack of features&#8217; effect. Because of this reason, it is used only from the largest single value as a non-linear feature. The combination of linear and non-linear characteristics as input is applied to MLP, Cascade Feed Forward and RBF neural networks and every (single) answer is studied. The answers for each class have a level of probability that any classifier can independently be taken to the classification of cardiac arrhythmias. A class that has the greatest probability is allocated to the data. These probabilities show the uncertainty of the answers. Each of the classifiers is considered as a witness. All the possibilities for different classes of each witness uncertainties function are modeled and crime function is defined. In other words, belief structure is formed for evidence. At this stage, by combined Demster law, the mass functions will combine together. In this situation, the level of uncertainty is much reduced and the class with the highest crime will be selected as the answer. According to the survey results, the combination of linear and non-linear characteristics for training and testing the neural networks classifiers has increased the accuracy of the answer. In other words, the extraction of more features leads to better training the neural networks and increases the accuracy of the classifiers. It can be noted that the using classifiers uncertainty principle and combining them by using the evidence theory has increased the accuracy of the final classification. The results of this study show that the proposed method was able to classify cardiac arrhythmias in the presence of noise and provided an acceptable answer for the intended issue. In sum, the proposed method has been able to classify 3 categories of cardiac arrhythmia such as PVC, PAC and NORMAL with high accuracy. This is performed in the best situation with sensitivity greater than 0/98. &#160;},  
Keywords = {ECG signal, Classifier, Neural Networks, Evidence theory},
volume = {14},
Number = {2}, 
pages = {25-42}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {طبقه‌بندی آریتمی‌های قلبی مبتنی بر ترکیب نتایج شبکه‌های عصبی با نظریه شواهد دمپستر- شفر},
abstract_fa ={آریتمی&#8204;های قلبی یکی از شایع&#8204;ترین &#160;بیماری&#8204;های قلبی است که ممکن است سبب مرگ بیمار شود. از&#8204;این&#8204;رو شناسایی آریتمی&#8204;های قلبی بسیار مهم است. در این مقاله برای دسته&#8204;بندی آریتمی&#8204;های قلبی در سه طبقه PAC، PVC و Normal روشی مبنی بر ترکیب طبقه&#8204;بندی&#8204;کننده&#8204;ها با استفاده از نظریه شواهد لحاظ شده است. بدین شکل که ابتدا پیک&#8204;های R در ECG شناسایی شد؛ سپس ویژگی&#8204;های&#160; خطی ECG شامل RMSSD، SDNN و HR Mean و همچنین ویژگی غیر خطی آن با استفاده از SVD به&#8204;دست آمد. ترکیب ویژگی&#8204;های &#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;&#160;به&#8204;دست&#8204;آمده به شبکه&#8204;های عصبی MLP، Cascade Feed Forward و RBF داده شد. اصل عدم قطعیت در مورد&#160; پاسخ آن&#8204;ها بررسی و در&#8204;نهایت پاسخ این طبقه&#8204;بندی&#8204;کننده&#8204;ها با استفاده از نظریه شواهد با یکدیگر ترکیب شدند. جهت پردازش ECG نیاز به حذف نوفه نبوده و روش پیشنهادی توانسته است در حضور نوفه، نوع آریتمی قلبی را در بهترین حالت با حساسیت 98 % تشخیص دهد. &#160;},
keywords_fa = {سیگنال الکتروکاردیوگرام, طبقه‌بندی‌کننده‌ها, شبکه‌های عصبی, نظریه شواهد },

doi = {10.18869/acadpub.jsdp.14.2.25},
url = {http://jsdp.rcisp.ac.ir/article-1-468-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-468-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Kianisarkaleh, Azadeh and Ghassemian, Mohammad Hass},  
title = {Modified Nonparametric Discriminant Analysis for Classification of Hyperspectral Images with Limited Training Samples}, 
abstract ={Feature extraction performs an important role in improving hyperspectral image classification. Compared with parametric methods, nonparametric feature extraction methods have better performance when classes have no normal distribution. Besides, these methods can extract more features than what parametric feature extraction methods do. Nonparametric feature extraction methods use nonparametric scatter matrices to compute transformation matrix. Nonparametric Discriminant Analysis (NDA) is one of the nonparametric feature extraction methods in which, to form nonparametric scatter matrices, local means of samples and weight function are used. Local mean is calculated by k nearest neighbors of each sample and weight function emphasizes on boundary samples in between class scatter matrix formation. In this paper, modified NDA (MNDA) is proposed to improve NDA. In MNDA, the number of neighboring samples, when measuring local mean, are determined considering position of each sample in feature space. MNDA uses new weight functions in scatter matrix formation. Suggested weight functions emphasizes on boundary samples in between class scatter matrix formation and focus on samples close to class mean in within class scatter matrix formation. Moreover, within class scatter matrix is regularized to avoid singularity. Experimental results on Indian Pines and Salinas images show that MNDA has better performance compared to other parametric, nonparametric feature extraction methods. For Indian Pines data set, the maximum average classification accuracy is 80.34%, which is obtained by 18 training samples, support vector machine (SVM) classifier and 10 extracted features achieved by MNDA method. For Salinas data set, the maximum average classification accuracy is 94.31%, which is obtained by 18 training samples, SVM classifier and 9 extracted features achieved by MNDA method. Experiments show that using suggested weight functions and regularized within class scatter matrix, the proposed method obtained better results in hyperspectral image classification with limited training samples. &#160;},  
Keywords = {Hyperspectral images, Feature extraction, Supervised classification, Hughes Phenomenon, Limited training samples},
volume = {14},
Number = {2}, 
pages = {43-58}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {تحلیل ممیز غیرپارامتریک بهبودیافته برای دسته‌بندی تصاویر ابرطیفی با نمونه آموزشی محدود},
abstract_fa ={استخراج ویژگی نقش مهمی در بهبود دسته&#173;بندی تصاویر ابرطیفی دارد. روش&#173;های استخراج ویژگی غیرپارامتریک، نسبت به روش&#173;های پارامتریک، برای داده&#173;های با توزیع غیر نرمال&#8204; کارایی بهتری دارند و می&#173;توانند ویژگی&#173;های بیشتری را استخراج کنند. روش&#173;های استخراج ویژگی غیرپارامتریک از ماتریس&#173;های پراکندگی غیرپارامتریک برای محاسبه ماتریس انتقال استفاده می&#173;کنند. تحلیل ممیز غیرپارامتریک[1]، یکی از روش&#173;های غیرپارامتریک در استخراج ویژگی است که در آن برای تشکیل ماتریس&#173;های پراکندگی غیرپارامتریک، از میانگین&#173;های محلی هر نمونه و تابع وزن استفاده می&#173;شود. میانگین محلی با استفاده از k نمونه همسایه به&#8204;دست می&#173;آید و تابع وزن، بر روی نمونه&#173;های مرزی در تشکیل ماتریس&#173; پراکندگی بین&#173;دسته&#173;ای تأکید می&#173;کند. در این مقاله، NDA بهبود&#8204;یافته[2] به&#8204;منظور اصلاح NDA معرفی شده است. در MNDA، تعداد نمونه&#173;های همسایه در محاسبه میانگین محلی با توجه به موقعیت نمونه در فضای ویژگی به&#8204;دست می&#173;آید. روش پیشنهادی از توابع وزن جدید در تشکیل ماتریس&#173;های پراکندگی استفاده می&#173;کند. توابع وزن پیشنهادی تأکید روی نمونه&#173;های مرزی در تشکیل ماتریس پراکندگی بین&#173;دسته&#173;ای و تأکید روی نمونه&#173;های نزدیک به میانگین دسته، در تشکیل ماتریس پراکندگی درون دسته&#173;ای دارند. علاوه براین، به&#8204;منظور اجتناب از تکین&#8204;شدن ماتریس پراکندگی درون&#8204;دسته&#173;ای، از تنظیم آن استفاده شده است. نتایج آزمایش&#173;ها روی تصاویر ایندیانا و سالیناس نشان می&#173;دهد که MNDA کاریی بهتری نسبت به روش&#173;های استخراج ویژگی پارامتریک و غیرپارامتریک مورد مقایسه داشته است. بیشترین مقدار صحت متوسط دسته&#173;بندی برای داده ایندیانا %34/80 است که با 18 نمونه آموزشی، دسته&#173;بند ماشین بردار پشتیبان و 10 ویژگی استخراج شده از MNDA به&#8204;دست آمده است. برای داده سالیناس، بیشترین مقدار صحت متوسط دسته&#173;بندی، %31/94 است که با 18 نمونه آموزشی، دسته&#173;بند ماشین بردار پشتیبان و 9 ویژگی استخراج&#8204;شده از MNDA به&#8204;دست آمده است. آزمایش&#173;ها نشان می&#8204;دهند که با استفاده از توابع وزن پیشنهادی و ماتریس پراکندگی درون&#8204;دسته&#173;ای تنظیم&#173;شده، روش پیشنهادی نتایج بهتری را در دسته&#8204;بندی تصاویر ابرطیفی با نمونه&#8204;های آموزشی محدود به&#8204;دست آورده است. &#160; [1] Nonparametric Discriminant Analysis (NDA) [2] Modified NDA (MNDA)},
keywords_fa = {تصاویر ابرطیفی, استخراج ویژگی, دسته‌بندی نظارت‌شده, پدیده هیوز, نمونه‌های آموزشی محدود},

doi = {10.18869/acadpub.jsdp.14.2.43},
url = {http://jsdp.rcisp.ac.ir/article-1-344-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-344-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Sajadi, Seyed mohamad bagher and Rashidi, Hassan and Minaeibidgoli, Behrooz},  
title = {A New Approach for Extracting Named Entity in Classical Arabic}, 
abstract ={In Natural Language Processing (NLP) studies, developing resources and tools makes a contribution to extension and effectiveness of researches in each language. In recent years, Arabic Named Entity Recognition (ANER) has been considered by NLP researchers due to a significant impact on improving other NLP tasks such as Machine translation, Information retrieval, question answering, query result clustering, etc. While most of these researches are based on Modern Standard Arabic (MSA), in this paper, we focus on Classical Arabic (CA) literature. We propose a corpus called NoorCorp with 130k labeled words for research purposes which is annotated by expert human resources manually. This corpus is based on a Historic-Islamic book of 1200 years ago including 1843 sentences and 127550 words. We also collected about 18k proper names from old Hadith books as a gazetteer which is called NoorGazet used as a future. In this paper, we propose a new approach to extract named entities (NEs) including person, location, organization and time. We use hybrid approach benefiting from advantages of Rule based approach and Machine learning approach. We divided the NoorCorp into two parts of training and test sets containing 80% and 20% of the data set respectively. Prediction model, based on Boosting method, was developed in two steps which Adaboost.M1 is employed to identify NEs and Adaboost.M2 is employed to classify NEs. There are many methods using multiple classifiers as voters and summing up their results, among which, ensemble methods are those which generate multiple hypotheses using the same base learner. We developed an ensemble consisting of 50 members (classifiers) based on decision stump to implement the weak learner. Since only 17% of the text data is composed of name entity labels, we had to deepen the tree while restricting pruning. We exploited tokenizing, part of speech (POS) tagging, and base phrase chunking (BPC) to overcome linguistic obstacles in Arabic including Meaning ambiguity, Optional diacritics, Complex morphology and Nonstandard written text. Moreover, using a statistical technique, the most frequently used words extracted as key words. Results show that performance of the method is better than decision tree as the base classifier. An overall F-measure value of 86.85 obtained which is better than base line about 20% and CART decision tree about 12%. Since CA corpus consists of simpler linguistic patterns compared to MSA, we applied the proposed approach on ANERCorp as Modern Standard Arabic corpus. Results show that the proposed model outcome on CA corpus is about 19% better than MSA. This result is due to the fact that there are plenty of NEs entered to MSA from other languages. These proper names do not have specific patterns and do not exist in the gazetteer. In addition, many NE&#8217;s are not distributed uniformly in ANERcorp which considerably reduces the results accuracy. &#160;&#160; &#160;},  
Keywords = {Named entity recognition (NER), Ensemble learning, Boosting method, Classical Arabic Language},
volume = {14},
Number = {2}, 
pages = {59-74}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {روشی جدید جهت استخراج موجودیت‌های اسمی در عربی کلاسیک},
abstract_fa ={تشخیص واحدهای اسمی به&#8204;عنوان یکی از سامانه&#8204;های پردازش زبان طبیعی عبارت از تشخیص اسامی خاص و طبقه&#8204;بندی آن&#8204;ها به یکی از گروه&#8204;های شخص، مکان، سازمان و زمان است. این عملیات به دلیل تأثیر قابل توجه در بهبود کارایی دیگر حوزه&#8204;های پردازش زبان طبیعی مانند ترجمه ماشین، بازیابی اطلاعات، خوشه&#8204;بندی نتایج جستجو و پرسش و پاسخ، در سال&#8204;های اخیر مورد توجه پژوهش&#8204;گران در زبان عربی نیز قرار گرفته است. گرچه بیشتر پژوهش&#8204;ها در این حوزه روی عربی استاندارد امروزی انجام &#8204;شده است، اما در این مطالعه عربی کلاسیک مورد توجه است. در همین راستا، روشی جدید جهت تشخیص واحدهای اسمی در زبان عربی ارائه می&#8204;شود. در این پژوهش یک پیکره متنی عربی کلاسیک به نام نورکورپ، متشکل از ۱۳۰ هزار کلمه برچسب&#8204;گذاری&#8204;شده توسط متخصصان، معرفی می&#8204;شود؛ همچنین از یک فرهنگ لغات شامل ۱۸۰۰۰ اسامی اشخاص که از کتب حدیثی استخراج شده است، به&#8204;عنوان منابع خارجی استفاده می&#8204;شود. مدل پیش&#8204;بینی، بر اساس مجمع رده&#8204;بندها و یک روش دو&#8204;مرحله&#8204;ای پیشنهاد شده است؛ به&#8204;طوری&#8204;که در مرحله نخست تشخیص واحدهای اسمی از طریق الگوریتم آدابوست M1 و در مرحله دوم طبقه&#8204;بندی آن&#8204;ها به گروه&#8204;های از&#8204;پیش&#8204;تعیین&#8204;شده توسط الگوریتم آدابوست M2 انجام می&#8204;شود. به&#8204;منظور غلبه بر چالش&#8204;های زبان عربی عملیات نشانه&#8204;گذاری، برچسب&#8204;گذاری ادات سخن و قطعه&#8204;کردن عبارت پایه به کار گرفته&#8204;شده است. با استفاده از یک روش آماری، برخی از کلمات پر کاربرد در واحدهای اسمی به&#8204;عنوان کلمات کلیدی استخراج شدند. نتیجه به&#8204;دست&#8204;آمده از مدل پیشنهادی در ارزیابی F-measure&#8204; معادل ۸۵/۸۶ درصد است که بیان&#8204;گر عملکرد مطلوب مدل است. در آخر، روش پیشنهادی روی یک پیکره استاندارد امروزی به نام انرکورپ اعمال و نتایج با پیکره نورکورپ مقایسه شده&#8204;اند. &#160;},
keywords_fa = {تشخیص واحدهای اسمی, مجمع رده‌بندها, روش بوستینگ, زبان عربی کلاسیک},

doi = {10.18869/acadpub.jsdp.14.2.59},
url = {http://jsdp.rcisp.ac.ir/article-1-295-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-295-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {SadatiTileboni, Seyyed Ali and Jazayeriy, Hamid and Valinataj, Mojtab},  
title = {Genetic Algorithm with Intelligence Chaotic Algorithm and Heuristic Multi-Point Crossover for Graph Coloring Problem}, 
abstract ={Graph coloring is a way of coloring the vertices of a graph such that no two adjacent&#160;vertices have the same color. Graph coloring problem (GCP) is about finding the smallest number of colors needed to color a given graph. The smallest number of colors needed to color a graph G, is called its chromatic number. GCP is a well-known NP-hard problems and, therefore, heuristic algorithms are usually used to solve it. GCP has many applications such as: bandwidth allocation, register allocation, VLSI design, scheduling, Sudoku, map coloring and so on. We try genetic algorithm (GA) and chaos theory to solve GCP. We proposed a heuristic algorithm called CMHn to implement multi-point crossover operation in GA. To generate initial population, a fast greedy algorithm is used. In this algorithm, the degree of each node and the number colors in its neighbor is used to assign a color to each node. Mutation operation in GA is used to explore the search space and scape from the local optima. In this study, a chaotic mutation operation is presented to select some vertices and change their color.&#160; The crossover and mutation parameters in the proposed algorithm is tuned based on some experiment. To evaluate the proposed algorithm, some experiment is conducted on DIMACS data set. Among DIMACS sample graphs, DSJ, Queen, Le450, Wap are well-known challenging samples for graph coloring. The proposed algorithm is executed 10 times on each sample and the best, worst and mean results are reported. Results show that the proposed algorithm can effectively solve GCP and have comparable outcome with the recent studies in this field. The proposed method outperforms other algorithms on very large graphs (Wap graphs).&#160; &#160;},  
Keywords = {Graph coloring problem, genetic algorithm, heuristics, multi-point crossover, chaotic mutation},
volume = {14},
Number = {2}, 
pages = {75-96}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {الگوریتم ژنتیک با جهش آشوبی هوشمند و ترکیب چند‌نقطه‌ای مکاشفه‌ای برای حل مسئله رنگ‌آمیزی گراف},
abstract_fa ={تخصیص مقدار رنگی را به هر یک از گره&#8204;های گراف، به&#8204;گونه&#8204;ای که هیچ دو گره مجاوری دارای رنگ یکسانی نباشد و کمترین مقدار رنگی استفاده شود، مسئله رنگ&#8204;آمیزی گراف گویند. این مسئله به&#8204;عنوان یکی از مسائل NP-hard شناخته می&#8204;شود که کاربردهای مختلفی در زمینه تخصیص پهنای باند، اختصاص حافظه به برنامه&#8204;ها و همچنین، طراحی مدارهای مجتمع دارد. در مقاله حاضر، از الگوریتم ژنتیک و پدیده آَشوب برای حل این مسئله استفاده شده است. در روش پیشنهادی حاضر، عمل&#8204;گر ترکیب چند&#8204;نقطه&#8204;ای مکاشفه&#8204;ای به نام CMHn معرفی شده است. این عمل&#8204;گر، با انتخاب چند نقطه برش در والدین و معتبر&#8204;کردن یکی از زیر بخش&#8204;های والدین (دومین زیربخش هر والد می&#8204;تواند معتبر یا غیر معتبر باشد) آنها را با هم، با استفاده از روشی ابتکاری ترکیب می&#8204;کند. برای اینکه بتوان از بهینه محلی فرار کرد و همچنین، برای یافتن فضای جستجوی جدید، از عمل&#8204;گر جهش استفاده می&#8204;شود. در این مقاله، عمل&#8204;گر جهش آشوبی هوشمند معرفی شده است که با استفاده از فرمولی گره&#8204;هایی را که برای جهش مناسب&#8204;ترند، انتخاب و بر روی آنها جهش را اعمال می&#8204;کند. همچنین، نیمی از جمعیت اولیه با استفاده از روش ابتکاری و نیمی از آن با روش تصادفی تولید شده است. به&#8204;منظور ارزیابی الگوریتم پیشنهادی از نمونه گراف&#8204;های DIMACS و Queen استفاده شده است. نتایج به&#8204;دست&#8204;آمده نشان می&#8204;دهد که روش پیشنهادی در بیش&#8204;تر گراف&#8204;ها، به&#8204;خصوص گراف&#8204;های بسیار بزرگ (wap) و گراف&#8204;های Queen جواب بهتری نسبت به تحقیقات مشابه ارائه می&#8204;دهد.},
keywords_fa = {مسئله رنگ‌آمیزی گراف, الگوریتم ژنتیک, روش ابتکاری, ترکیب چند نقطه‌ای مکاشفه‌ای, جهش آشوبی هوشمند},

doi = {10.18869/acadpub.jsdp.14.2.75},
url = {http://jsdp.rcisp.ac.ir/article-1-392-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-392-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Jamshidi, Ali and Yazdi, Mehran and Manafi, Maryam},  
title = {Image Compression Based on Intelligent Information Removing and Inpainting Reconstruction Algorithms}, 
abstract ={Compression can be done by lossy or lossless methods. The lossy methods have been used more widely than the lossless compression. Although, many methods for image compression have been proposed yet, the methods using intelligent skipping proper to the visual models has not been considered in the literature. Image inpainting refers to the application of sophisticated algorithms to replace lost or corrupted parts of the data so that visual difference cannot be inferred from the reconstructed image. In this paper, first we review some of the image inpainting algorithms and some of the image compression techniques using the inpainting algorithms, we propose a new inpainting based image compression algorithm that can improve the compression rate considerably. We present image compression system based on the proposed parameter-assistant image inpainting method to more deeply exploit visual redundancy inherent in color images. We have shown that with carefully selected dropped regions and appropriately extracted parameters from them, dropped regions can be satisfactorily restored using the proposed PAI algorithm. Accordingly, our compression scheme has a higher coding performance compared with traditional methods in terms of the perceptual quality. To best represent the target region for inpainting, an effective region classifier is required. A generic solution is to study the distribution of each image region and find the best match among the candidates in the predefined model class. For simplicity, in our scheme, an entire image divided into three categories: gradated, structural, and non-featured, at non-overlapping block level of size S&#215;S. The classification is performed based on edge content and color variance in each block. Simulation results show that our proposed method has reasonable visual quality in comparison with the other proposed image compression algorithms.&#160;&#160; &#160;},  
Keywords = {Image Inpainting, Image Compression, Intelligent Information Removing, Coding },
volume = {14},
Number = {2}, 
pages = {97-114}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {فشرده‌سازی تصویر با کمک حذف و کدگذاری هوشمندانه اطلاعات تصویر و بازسازی آن با استفاده از الگوریتم های ترمیم تصویر},
abstract_fa ={روش&#8204;های فشرده&#8204;سازی با اتلاف، به&#8204;دلیل ایجاد فشرده&#8204;سازی بیشتر، کاربرد گسترده&#8204;تری دارند. اگرچه روش&#8204;های زیادی تا به حال برای فشرده&#8204;سازی تصاویر پیشنهاد شده ، اما،&#160; به استفاده از روش&#8204;های هوشمندانه حذف اطلاعات، کمتر توجه شده است. ترمیم، مجموعه&#8204;ای از روش&#8204;هایی است که اصلاحاتی را بر روی تصاویر انجام می&#8204;دهد؛ با این هدف که بیننده تفاوتی بین تصویر اصلاح&#8204;شده و تصویر اصلی احساس نکند. در این مقاله، پس از بررسی و معرفی بعضی روش&#8204;های ترمیم تصویر و روش&#8204;های فشرده&#8204;سازی تصویر با کمک ترمیم، روش جدیدی پیشنهاد می&#8204;شود که علاوه&#8204;بر&#8204;این که باعث فشردگی قابل توجه تصویر در زمان ارسال می&#8204;شود، نتیجه کیفی مناسبی نیز در گیرنده خواهد داشت. در روش پشنهادی، تصویر به نواحی ساختاری و بافتی تقسیم می&#8204;شود و برای هر ناحیه بلوک&#8204;های قابل حذفی که امکان بازسازی مناسبی در گیرنده با استفاده از روش&#8204;های ترمیم دارند، شناسایی و حذف می&#8204;شوند و اطلاعات کمکی لازم جهت ترمیم بهتر از آنها استخراج می&#8204;شود. این بلوک&#8204;ها به&#8204;همراه بلوک&#8204;های غیرقابل حذف تصویر پس از کد&#8204;شدن، ارسال می&#8204;شوند و در گیرنده پس از کدگشایی، بلوک&#8204;های از&#8204;دست&#8204;رفته بازسازی و ترمیم می&#8204;گردند تا در&#8204;نهایت تصویر اولیه در گیرنده قابل استفاده باشد. ویژگی&#8204;های روش پیشنهادی نخست متغیر&#8204;بودن اندازه بلوک&#8204;های حذفی است که باعث فشردگی بیشتر می&#8204;شود و ثانیاً ارائه روش جدیدی جهت بازسازی بلوک&#8204;های شامل لبه در گیرنده است که کیفیت بلوک&#8204;های ترمیم شده این نواحی را افزایش می&#8204;دهد. &#160;},
keywords_fa = {ترمیم تصویر, فشرده‌سازی تصاویر, حذف هوشمندانه اطلاعات, کد کردن },

doi = {10.18869/acadpub.jsdp.14.2.97},
url = {http://jsdp.rcisp.ac.ir/article-1-434-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-434-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {},  
title = {A New Method for Classification of Nano-Structures based on Time Series Analysis and Fuzzy Logic}, 
abstract ={Dispersion of nanoparticles in nanostructures is one of the most important indicators designed to verify the effectiveness of proposed methods in the synthesis of nanomaterials. In the recent years, various methods have been suggested for the synthesis of nanostructures in which the Scanning Electron Microscopy (SEM) has been used to show the quality of the nanomaterial. The SEM images of nanoparticles contain structural, chemical and morphological information with high resolution in nanometer scale of nanomaterials. One of the challenges in the quality of dispersion&#8217;s nanostructures is detection of agglomeration degree. In some SEM images of nanoparticles, the particles have speeded uniformly and not aggregately. In some of the other SEM images, their particles are agglomerated. Also, there are a few SEM images of nanoparticles that their particles aren&#8217;t very aggregate or diffused. If the SEM images of nanoparticles with their particles speeded uniformly, are called good images, and the images with their aggregate particles are called bad images, and the images with their particle dispersion between good and bad images, are called average images, the nanomaterials could be classified in categories of good, average, and bad images. In this paper, a new algorithm has been provided to classify nanostructures using SEM images of nanoparticles. For this purpose, these images were transformed to time series at first (the time series extracted are unique for each SEM image of nanoparticles) and their specifications were investigated through time series analysis methods. Then, statistical specifications of these series were extracted. Six statistical specifications have been extracted for classification of nanostructures. These specifications are as follows: standard deviation, first and second kurtosis, interquartile range, the criterion of Pearson, and skewness. The extracted specifications were used as inputs of a fuzzy inference system for classifying microscopic images of nanostructures into three groups: good, average and bad. This algorithm has been tested on 65 nanoparticles microscopic images with identical size and resulted precision above 93 percent indicated validity of this algorithm. &#160;},  
Keywords = {SEM image of nanoparticles, Time series analyses, Statistical features, Fuzzy logic},
volume = {14},
Number = {2}, 
pages = {115-130}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {یک روش جدید برای طبقه‎بندی نانوساختارها براساس آنالیز سری زمانی و منطق فازی},
abstract_fa ={میزان پراکندگی نانوذرات در نانوساختارها، از مهم&#8204;ترین شاخص&#8204;هایی است که جهت تأیید کارآیی روش&#8204;های پیشنهادی در زمینه سنتز نانومواد به&#8204;کار می&#8206;رود. تصاویر میکروسکوپی الکترونی روبشی نانوذرات دارای اطلاعات ساختاری، شیمیایی و مورفولوژیکی با وضوح بالا در مقیاس نانومتری نانومواد هستند. در این مقاله، یک الگوریتم جدید جهت طبقه&#8206;بندی نانوساختارها با استفاده از این تصاویر ارائه شده &#8206;است؛ بدین منظور، ابتدا تصاویر میکروسکوپی الکترونی روبشی نانوذرات به سری زمانی تبدیل و مشخصات آنها از طریق روش&#8204;های تحلیل سری زمانی مورد بررسی قرار گرفتند؛ سپس ویژگی&#8204;های آماری این سری&#8204;ها استخراج و به&#8204;عنوان ورودی&#8206;های یک سامانه استنتاج فازی برای طبقه&#8206;بندی تصاویر میکروسکوپی نانوساختارها در سه گروه خوب، متوسط و بد در نظر گرفته &#8206;شدند. این الگوریتم برروی 65 تصویر میکروسکوپی نانوذرات با ابعاد یکسان (250&#215;250 پیکسل) اعمال شده و دقتی بالاتر از 93 درصد را به دنبال داشته &#8206;است که بسیار مناسب است. &#160;},
keywords_fa = {تصاویر میکروسکوپی الکترونی روبشی نانوذرات, تحلیل سری زمانی, ویژگی‎های آماری, منطق فازی},

doi = {10.18869/acadpub.jsdp.14.2.115},
url = {http://jsdp.rcisp.ac.ir/article-1-409-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-409-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Mikaeili, Mohammad and Najafi, foroogh},  
title = {Performance Analysis of a Persian text input brain–computer interface (BCI) P300 Speller system with row/column paradigm (RCP)}, 
abstract ={As a Brain computer interface system, BCI P300 Speller tries to help disabled people and patients to regain some of their lost ability with allowing communication via typing. The ability of personalization is one of the most important features in a BCI system, so the typing language as a personalization factor is an important feature in a BCI speller. Most prior researches on P300 Speller has focused on displaying English alphabet and there were only few studies made on other languages such as Chinese. In this research, we present a P300 Speller system, based on RCP, for Persian (Farsi) character input. RCP (Row or Column Paradigm) was introduced by Farwell and Donchin at 1988, and since then it has been considered as a benchmark in P300 BCI speller research. As a result, in this study also, Row or column paradigm was selected as the base stimulation pattern in P300 speller system. In order to evaluate the Persian row or column paradigm performance, we recorded EEG signals from volunteered subjects while the stimulation pattern was being displayed. It should be noted that the test was explained to each subject before testing, and for more experience and in order to reduce the error, each subject participated in an experiment test before attending the main test. These EEG signals were recorded from 8 channels based on &#8216;&#8216;Fz&#8217;&#8217;, &#8216;&#8216;Cz&#8217;&#8217;, &#8216;&#8216;P3&#8217;&#8217;, &#8216;&#8216;Pz&#8217;&#8217;, &#8216;&#8216;P4&#8217;&#8217;, &#8216;&#8216;O1&#8217;&#8217;, &#8216;&#8216;Oz&#8217;&#8217; and &#8216;&#8216;O2&#8217;&#8217; site in accordance to the International 10&#8211;20 system electrode placement system and by using Science Beam co.&#8217;s EEG recording device. The sample rate was 1 KHz which was down sampled to 250Hz. After recording, the EEG signals were filtered using a band passed filter And for classification, Linear discriminate analysis was used in combination with K-fold validation method for classifier training. As performance determination, we calculated accuracy and bit rate for the mentioned system based on recorded data from volunteers and reached the average accuracy of 88.21% and bit rate of 6.74 (bits/minute) (we use Linear LDA classifier for classification and the total trial number was set to 15). Furthermore, in this research performance was measured for different trial number and final results demonstrated that this system can achieve high average accuracy of 80.06% and average bit rate of 42.43 (bits/minute) by using only 2 repetitions. &#160;},  
Keywords = {Brian computer interface systems, BCI P300 Speller, P300 wave, LDA classifier},
volume = {14},
Number = {2}, 
pages = {131-140}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {سنجش عملکرد سامانه‌های رابط مغز و رایانه P300 Speller به‌ازای ماتریس نمایش ردیف و یا ستون  (RCP) و نمایش حروف زبان فارسی},
abstract_fa ={سامانه&#8204;های رابط مغز و رایانه P300 Speller به&#8204;عنوان عضوی از خانواده سامانه&#8204;های رابط مغز و رایانه سعی دارد تا توانایی تایپ حروف و برقراری ارتباط با این شیوه را برای بیماران و معلولان فراهم آورد. یکی از موارد بسیار مهم در این سامانه&#8204;ها، قابلیت شخصی&#8204;سازی است. با توجه به آنکه بیش&#8204;تر پژوهش&#8204;های این حوزه بر اساس نمایش حروف انگلیسی انجام شده، در این پژوهش سعی شده است تا برای نخستین&#8204;بار عملکرد یک سامانه ارتباط مغز و رایانه P300 Speller به&#8204;ازای نمایش حروف زبان فارسی مورد سنجش قرار گیرد. در این پژوهش پس از ثبت داده از داوطلبان و سنجش عملکرد سامانه مورد بررسی، صحت تشخیص 21/88% و نرخ انتقال اطلاعات 74/6 بیت در دقیقه به&#8204;ازای پانزده تکرار با &#160;ترکیب روش کاهش بعد LDA و طبقه بند بیز به&#8204;دست آمد. همچنین در این پژوهش اثر تغییر تعداد تکرار و کاهش زمان آزمایش نیز مورد بررسی قرار گرفت و نشان داده شد که به&#8204;ازای کمینه تعداد تکرار، می&#8204;توان به صحت تشخیص 06/80% و نرخ انتقال اطلاعات 43/42 بیت بر دقیقه دست یافت. &#160;},
keywords_fa = {سامانه‌های رابط مغز و رایانه, سامانه P300 Speller BCI, روش کاهش بعد LDA, مؤلفه P300 , طبقه‌بند بیز},

doi = {10.18869/acadpub.jsdp.14.2.131},
url = {http://jsdp.rcisp.ac.ir/article-1-448-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-448-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {TolouBeidokhti, Mohammad Amin and Ahmadyfard, Alirez},  
title = {Document Image Dewarping using geometrical information extracted from document lines}, 
abstract ={Document images produced by scanners or digital cameras usually have photometric and geometric distortions. If either of these effects distorts document, recognition of words from such a document image using OCR is subject to errors. In this paper we propose a novel approach to significantly remove geometric distortion from document images. In this method first we extract document lines from document using morphological operators. Then, extracted document lines are divided into a number of equal size column strips.&#160; This allows to assume that each segment of line document is not curved. Each extracted document line segment is aligned horizontally. For this purpose, a segment line of document is rotated at different angels and for each rotation horizontal projection is obtained. The rotation angle with maximum peak at the corresponding projection signal is selected to align the line segment, horizontally. In order to estimate the geometrical distortion, for each document line a reference point is extracted from each line segment. These points indicate the position of a document line at starting column of line segments. Using reference points of a document line a polynomial function is fitted to each document line. At the end, geometric distortion for each part of the document is eliminated using a perspective transformation. This transformation is estimated based on the extracted polynomial function. To increase the stability of the proposed method for short text lines, the curve of adjacent text lines of longer length is used. A post processing stage is required after applying perspective transformation on document patches. Since this transformation is a continuous mapping but it is applied on digital images. To remove this distortion from the result, the consistency of each pixel value with the value of neighboring pixels are considered to correct the value of inconsistence pixels. The proposed method is implemented on Persian and English databases and has been compared with the existing methods. The results indicate the efficiency and accuracy of the proposed method in elimination of geometric distortions. &#160;},  
Keywords = {Geometric distortion, document processing, perspective Transformation, Optical character recognition (OCR)},
volume = {14},
Number = {2}, 
pages = {141-158}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {رفع اعوجاج هندسی متون به‌کمک
 اطلاعات هندسی خطوط متن},
abstract_fa ={تصاویر سند تهیه&#8204;شده توسط پویش&#8204;گر یا دوربین دیجیتال، همواره با اعوجاج&#8204;های فتومتریک و هندسی همراه هستند. وجود هر دو نوع اعوجاج، باعث کاهش دقت عملکرد نرم&#8204;افزارهای شناسایی نویسه&#173;ها (OCR) می&#173;شوند. در این مقاله روشی نوین جهت رفع اعوجاج&#8204;های هندسی از تصاویر متنی ارائه شده &#173;است. در روش پیشنهادی به&#8204;منظور تصحیح اعوجاج هندسی، در ابتدا خطوط متن از تصویر استخراج و سپس هر خط متن به ستون&#173;هایی به عرض مساوی شکسته می&#173;شوند. برای هر قطعه استخراج&#8204;شده از یک خط، راستای قطعه به&#8204;نحوی تصحیح می&#8204;شود که حروف موجود در آن قطعه در راستای افقی قرار گیرد. برای این منظور به&#8204;ازای چرخش&#173;های مختلف قطعۀ متن، افکنش افقی تصویر محاسبه می&#173;شود و چرخشی از قطعه که بلندترین قله افکنش را ایجاد کند، راستای تصحیح&#8204;شده آن قطعه در نظر گرفته می&#8204;شود. بر این اساس یک نقطه مرجع که معرف راستای مبنا است، برای هر قطعه&#173;خط هم&#8204;راستا&#173;شده با افق استخراج می&#8204;شود. به&#8204;کمک نقاط مرجع، هر قطعه از خط، انحنای آن خط متن به&#8204;کمک برازش یک تابع درجۀ سه به&#8204;دست می&#8204;آید. درنهایت با استفاده از تخمین تبدیل پرسپکتیو، اعوجاج هندسی هر خط برطرف می&#8204;شود. جهت افزایش پایداری روش پیشنهادی در تخمین انحنای خطوط متن با طول کم، از انحنای خطوط با طول بزرگ&#173;تر مجاور آن خط استفاده شده &#8204;است. روش&#173; پیشنهادی بر روی پایگاه&#173;های دادۀ فارسی و انگلیسی پیاده&#173;سازی و با برخی روش&#8204;های هم&#8204;تراز آن مقایسه شده است. نتایج بیان&#8204;گر قدرت و دقّت روش پیشنهادی در رفع اعوجاج هندسی است. &#160;},
keywords_fa = {اعوجاج هندسی, پردازش دوبعدی اسناد, تخمین تبدیل پرسپکتیو, نویسه‌خوان نوری},

doi = {10.18869/acadpub.jsdp.14.2.141},
url = {http://jsdp.rcisp.ac.ir/article-1-406-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-406-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

@article{ 
author = {Chaghari, Arash and Feizi-Derakhshi, Mohammad-Rez},  
title = {Automatic Clustering Using Improved Imperialist Competitive Algorithm}, 
abstract ={Imperialist Competitive Algorithm (ICA) is considered as a prime meta-heuristic algorithm to find the general optimal solution in optimization problems. This paper presents a use of ICA for automatic clustering of huge unlabeled data sets. By using proper structure for each of the chromosomes and the ICA, at run time, the suggested method (ACICA) finds the optimum number of clusters while optimal clustering of the data simultaneously.To increase the accuracy and speed of convergence, the structure of ICA changes. As in different applications, there is a need for data clustering which the number of clusters is not known before it is necessary to have methods that can cluster data without knowing the correct prediction of the number of clusters. In the other words, the proposed algorithm requires no background knowledge to classify the data.&#160; In addition, the proposed method is more accurate in comparison with other clustering methods based on evolutionary algorithms. In Imperialist Competitive Algorithm, firstly steps should be taken to increase search rates and explore possible solution while approaching to the global optimal response the steps should be reduced to ensure that the algorithm is not lost and it is not in the local optimal manner. For this purpose and improvement of imperialist competitive algorithm, mutation rate and revolution operator&#39;s operation rate are determined dynamically. DB and CS are cluster validity Indexes. In this paper, DB and CS cluster validity measurements are used as the objective function. To demonstrate the superiority of the proposed method, the average of fitness function and the number of clusters determined by the proposed method is compared with three automatic clustering algorithms based on evolutionary algorithms. The partitional clustering algorithms are based on three powerful well-known optimization algorithms, namely the genetic algorithm, the particle swarm optimization and differential evolutionary algorithm.},  
Keywords = {Partitional Clustering, Automatic Clustering, Imperialist Competitive Algorithm (ICA)},
volume = {14},
Number = {2}, 
pages = {159-169}, 
publisher = {Research Center on Developing Advanced Technologies},
title_fa = {خوشه‌بندی خودکار داده‌ها با بهره‌گیری از الگوریتم رقابت استعماری بهبودیافته},
abstract_fa ={الگوریتم رقابت استعماری (ICA)، یکی از کاراترین الگوریتم&#8204;های فرا&#8204;ابتکاری برای پیدا&#8204;کردن جواب بهینه سراسری در مسائل بهینه&#8204;سازی است. در این مقاله از الگوریتم رقابت استعماری برای خوشه&#8204;بندی خودکار مجموعه داده&#8204;های بزرگ و واقعی بدون برچسب استفاده شده است. با بهره&#8204;گیری از ساختار مناسب برای هر یک از کروموزم&#8204;ها و استفاده از الگوریتم رقابت استعماری، در زمان اجرا تعداد بهینه خوشه&#8204;ها هم&#8204;زمان با خوشه&#8204;بندی بهینه داده&#8204;ها به&#8204;دست می&#8204;آید. همچنین برای افزایش دقت و افزایش سرعت هم&#8204;گرایی، ساختار الگوریتم رقابت استعماری با تغییراتی همراه است. روش پیشنهادی (ACICA) نیاز به هیچ&#8204;گونه دانش قبلی برای خوشه&#8204;بندی داده&#8204;ها ندارد. علاوه&#8204;بر آن روش پیشنهادی&#160; در مقایسه با سایر روش&#8204;های خوشه&#8204;بندی مبتنی بر الگوریتم&#8204;های تکاملی، دقت بیشتری را دارد. از معیارهای ارزیابی خوشه&#8204;بندی DB و CS به&#8204;عنوان تابع هدف استفاده شده است. برای نشان&#8204;دادن برتری روش پیشنهادی، میانگین مقدار بهینه تابع هدف و تعداد خوشه&#173;های تعیین&#8204;شده توسط روش پیشنهادی با سه الگوریتم خوشه&#173;بندی خودکار مبتنی بر الگوریتم&#173;های تکاملی مقایسه می&#8204;شود. &#160;},
keywords_fa = {خوشه‌بندی تفکیکی, خوشه‌بندی خودکار, الگوریتم رقابت استعماری},

doi = {10.18869/acadpub.jsdp.14.2.159},
url = {http://jsdp.rcisp.ac.ir/article-1-453-en.html},  
eprint = {http://jsdp.rcisp.ac.ir/article-1-453-en.pdf},  
journal = {Signal and Data Processing},  
issn = {2538-4201}, 
eissn = {2538-421X}, 
year = {2017}  
}

