refactor project structure

This commit is contained in:
2021-08-07 01:16:56 +02:00
parent bf30511678
commit 53fdb7530b
23 changed files with 22 additions and 19 deletions
+4 -1
View File
@@ -36,9 +36,12 @@ fetch_batch <- function(offset, download_dir) {
#' This fetches all available records of the 19th legislative period of the german Bundestag.
#'
#' @param download_dir character
#' @param create bool
#'
#' if create is TRUE, the directory given in download_dir is created
#'
#' @export
fetch_all <- function(download_dir="records/", create=FALSE) {
fetch_all <- function(download_dir="data/records/", create=FALSE) {
# check if download_dir path is a directory path
if (str_sub(download_dir, -1) != .Platform$file.sep)
download_dir <- str_c(download_dir, .Platform$file.sep)
+3 -3
View File
@@ -8,7 +8,7 @@
#' @param path character
#'
#' @export
read_all <- function(path="records/") {
read_all <- function(path="data/records/") {
cat("Reading all records from", path, "\n")
available_protocols <- list.files(path)
res <- pblapply(available_protocols, read_one, path=path)
@@ -214,7 +214,7 @@ parse_speakerlist <- function(speakerliste_xml) {
#' if create is set to TRUE, the directory given in path is created
#'
#' @export
write_to_csv <- function(tables, path="csv/", create=F) {
write_to_csv <- function(tables, path="data/csv/", create=F) {
check_directory(path, create)
write.table(tables$speaker, str_c(path, "speaker.csv"))
write.table(tables$speeches, str_c(path, "speeches.csv"))
@@ -230,7 +230,7 @@ write_to_csv <- function(tables, path="csv/", create=F) {
#' Reading the tables from a csv is way faster than reading and repairing the data every single time
#'
#' @export
read_from_csv <- function(path="csv/") {
read_from_csv <- function(path="data/csv/") {
list(speaker = read.table(str_c(path, "speaker.csv")) %>%
tibble() %>%
mutate(id = as.character(id)),