from arabica import coffee_break coffee_break(text = data['text'], time = data['period'], date_format = 'us', # Use US-style date format to read dates preprocess = True, # Clean data - digits and punctuation skip = ['
', # Remove additional stop words '/n', 'Another Long String'], n_breaks = 2, # 2 breakpoints identified time_freq = 'Y') # Yearly aggregation