<?xml version="1.0" encoding="utf-8" standalone="yes"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
	<channel>
		<title>Plan on Dataprd.Com</title>
		<link>https://dataprd.com/tags/plan/</link>
		<description>Recent content in Plan on Dataprd.Com</description>
		<generator>Hugo</generator>
		<language>en-us</language>
		
		
		
		
			<lastBuildDate>Wed, 12 Mar 2014 18:17:16 +0000</lastBuildDate>
		
			<atom:link href="https://dataprd.com/tags/plan/index.xml" rel="self" type="application/rss+xml" />
			<item>
				<title>A multi-tiered Big Data warehouse &amp; processing facility</title>
				<link>https://dataprd.com/posts/a-multi-tiered-big-data-warehouse-processing-facility/</link>
				<pubDate>Wed, 12 Mar 2014 18:17:16 +0000</pubDate>
				<guid>https://dataprd.com/posts/a-multi-tiered-big-data-warehouse-processing-facility/</guid>
				<description>&lt;p&gt;The article details an exemplary setup for a multi-tiered data warehouse and processing facility using Hadoop and batch type of data analysis. Requirements in brief:&lt;/p&gt;&#xA;&lt;ul&gt;&#xA;&lt;li&gt;High peaks on data acquisition (e.g. financial transactions of yearly festivals, R&amp;amp;D facility with images or videos flowing in only at the time of experimenting)&#xA;&lt;ul&gt;&#xA;&lt;li&gt;Velocity is high at peak times, Gigabytes / second for short time (max. few hours / day)&lt;/li&gt;&#xA;&lt;li&gt;Data to be quickly accessed for data processing for some weeks after acquisition&lt;/li&gt;&#xA;&lt;/ul&gt;&#xA;&lt;/li&gt;&#xA;&lt;li&gt;High volume of data to be stored for long-time&#xA;&lt;ul&gt;&#xA;&lt;li&gt;Data to be accessed some years after acquisition, access can be delayed by days after request&lt;/li&gt;&#xA;&lt;/ul&gt;&#xA;&lt;/li&gt;&#xA;&lt;/ul&gt;&#xA;&lt;p&gt;The below setup details 3 tiers:&lt;/p&gt;</description>
			</item>
	</channel>
</rss>
