{"id":11253,"date":"2022-05-18T10:00:06","date_gmt":"2022-05-18T08:00:06","guid":{"rendered":"http:\/\/sidor.webs5.uvigo.es\/?p=11253"},"modified":"2023-03-07T10:25:08","modified_gmt":"2023-03-07T09:25:08","slug":"2022-bdpar-big-data-preprocessing-architecture","status":"publish","type":"post","link":"https:\/\/sidor.uvigo.es\/en\/2022-bdpar-big-data-preprocessing-architecture\/","title":{"rendered":"2022 \/\/ bdpar: Big Data Preprocessing Architecture"},"content":{"rendered":"<p><b>Authors: Miguel Ferreiro-D\u00edaz [aut, cre], David Ruano-Ord\u00e1s [aut, ctr], Tom\u00e1s R. Cotos-Ya\u00f1ez [aut, ctr], Jos\u00e9 Ram\u00f3n M\u00e9ndez Reboredo [aut, ctr], University of Vigo [cph]<\/b><\/p>\n<h3>Description<\/h3>\n<hr class=\"hr\" \/>\n<p>Provide a tool to easily build customized data flows to pre-process large volumes of information from different sources. To this end, &#8216;bdpar&#8217; allows to (i) easily use and create new functionalities and (ii) develop new data source extractors according to the user needs. Additionally, the package provides by default a predefined data flow to extract and pre-process the most relevant information (tokens, dates, &#8230; ) from some textual sources (SMS, Email, tweets, YouTube comments).<\/p>\n<p>&nbsp;<\/p>\n<p><b>R Package \u2013 Version 3.0.2<\/b><br \/>\n<b>Published:<\/b>\u00a02022-05-18<br \/>\n<b>Link:<\/b>\u00a0<a href=\"https:\/\/cran.r-project.org\/web\/packages\/bdpar\/index.html\" target=\"_blank\" rel=\"noopener noreferrer\">https:\/\/cran.r-project.org\/web\/packages\/bdpar\/index.html<\/a><\/p>\n<p><iframe loading=\"lazy\" src=\"https:\/\/cranlogs.r-pkg.org\/badges\/bdpar\" width=\"300\" height=\"50\"><\/iframe><\/p>\n","protected":false},"excerpt":{"rendered":"<p>Authors: Miguel Ferreiro-D\u00edaz [aut, cre], David Ruano-Ord\u00e1s [aut, ctr], Tom\u00e1s R. Cotos-Ya\u00f1ez [aut, ctr], Jos\u00e9 Ram\u00f3n M\u00e9ndez Reboredo [aut, ctr], University of Vigo [cph] Description &hellip; <a class=\"excerpt_more\" href=\"https:\/\/sidor.uvigo.es\/en\/2022-bdpar-big-data-preprocessing-architecture\/\">Continued<\/a><\/p>\n","protected":false},"author":2,"featured_media":0,"comment_status":"closed","ping_status":"closed","sticky":false,"template":"","format":"standard","meta":{"footnotes":""},"categories":[107],"tags":[122],"class_list":["post-11253","post","type-post","status-publish","format-standard","hentry","category-software-development","tag-software"],"jetpack_featured_media_url":"","_links":{"self":[{"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/posts\/11253\/"}],"collection":[{"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/posts\/"}],"about":[{"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/types\/post\/"}],"author":[{"embeddable":true,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/users\/2\/"}],"replies":[{"embeddable":true,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/comments\/?post=11253"}],"version-history":[{"count":1,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/posts\/11253\/revisions\/"}],"predecessor-version":[{"id":11254,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/posts\/11253\/revisions\/11254\/"}],"wp:attachment":[{"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/media\/?parent=11253"}],"wp:term":[{"taxonomy":"category","embeddable":true,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/categories\/?post=11253"},{"taxonomy":"post_tag","embeddable":true,"href":"https:\/\/sidor.uvigo.es\/en\/wp-json\/wp\/v2\/tags\/?post=11253"}],"curies":[{"name":"wp","href":"https:\/\/api.w.org\/{rel}","templated":true}]}}