<?xml version="1.0" encoding="UTF-8"?>
<rss xmlns:content="http://purl.org/rss/1.0/modules/content/" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:rdf="http://www.w3.org/1999/02/22-rdf-syntax-ns#" xmlns:taxo="http://purl.org/rss/1.0/modules/taxonomy/" version="2.0">
  <channel>
    <title>topic extract data into OneLake using Fabric in Pipelines</title>
    <link>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041129#M4547</link>
    <description>&lt;P&gt;I have a collection of PDFs from various sources and need to extract data into OneLake using Fabric.&lt;BR /&gt;How can I do that? Are there any documentation or steps I should follow?&lt;/P&gt;</description>
    <pubDate>Sun, 14 Jul 2024 14:50:49 GMT</pubDate>
    <dc:creator>esraaE</dc:creator>
    <dc:date>2024-07-14T14:50:49Z</dc:date>
    <item>
      <title>extract data into OneLake using Fabric</title>
      <link>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041129#M4547</link>
      <description>&lt;P&gt;I have a collection of PDFs from various sources and need to extract data into OneLake using Fabric.&lt;BR /&gt;How can I do that? Are there any documentation or steps I should follow?&lt;/P&gt;</description>
      <pubDate>Sun, 14 Jul 2024 14:50:49 GMT</pubDate>
      <guid>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041129#M4547</guid>
      <dc:creator>esraaE</dc:creator>
      <dc:date>2024-07-14T14:50:49Z</dc:date>
    </item>
    <item>
      <title>Re: extract data into OneLake using Fabric</title>
      <link>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041340#M4549</link>
      <description>&lt;P&gt;Hi&amp;nbsp;&lt;a href="javascript:void(0)" data-lia-user-mentions="" data-lia-user-uid="747283" data-lia-user-login="esraaE" class="lia-mention lia-mention-user"&gt;esraaE&lt;/a&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;First extract the data, you might consider using a tool or library to extract the data from the PDF.&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;After extracting the data, you may need to clean the data and convert it into a suitable format (for example, CSV, JSON) to load into OneLake.&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;Here are some information about connecting to onelake:&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;&lt;A href="https://learn.microsoft.com/en-us/fabric/onelake/onelake-access-api" target="_blank"&gt;How do I connect to OneLake? - Microsoft Fabric | Microsoft Learn&lt;/A&gt;&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;&lt;A href="https://learn.microsoft.com/en-us/fabric/data-engineering/load-data-lakehouse" target="_blank"&gt;Options to get data into the Lakehouse - Microsoft Fabric | Microsoft Learn&lt;/A&gt;&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;&lt;A href="https://learn.microsoft.com/en-us/fabric/onelake/onelake-access-python" target="_blank"&gt;Access OneLake with Python - Microsoft Fabric | Microsoft Learn&lt;/A&gt;&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;Regards,&lt;/P&gt;
&lt;P&gt;Nono Chen&lt;/P&gt;
&lt;P&gt;If this &lt;STRONG&gt;&lt;EM&gt;post&lt;/EM&gt;&lt;/STRONG&gt;&amp;nbsp;helps, then please consider&amp;nbsp;&lt;STRONG&gt;&lt;EM&gt;Accept it as the solution&lt;/EM&gt;&lt;/STRONG&gt;&amp;nbsp;to help the other members find it more quickly.&lt;/P&gt;</description>
      <pubDate>Mon, 15 Jul 2024 01:59:16 GMT</pubDate>
      <guid>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041340#M4549</guid>
      <dc:creator>Anonymous</dc:creator>
      <dc:date>2024-07-15T01:59:16Z</dc:date>
    </item>
    <item>
      <title>Re: extract data into OneLake using Fabric</title>
      <link>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041463#M4550</link>
      <description>&lt;P&gt;Assuming the sources are supported by either Dataflow Gen 2 or datapipelines, you can extract data from PDF via dataflow Gen 2 .&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;Below blog explains the same:&lt;/P&gt;
&lt;P&gt;&lt;A href="https://datasharkx.wordpress.com/2023/12/03/read-and-import-data-from-pdf-file-using-msft-fabric/" target="_blank"&gt;https://datasharkx.wordpress.com/2023/12/03/read-and-import-data-from-pdf-file-using-msft-fabric/&lt;/A&gt;&lt;/P&gt;</description>
      <pubDate>Mon, 15 Jul 2024 03:13:59 GMT</pubDate>
      <guid>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4041463#M4550</guid>
      <dc:creator>NandanHegde</dc:creator>
      <dc:date>2024-07-15T03:13:59Z</dc:date>
    </item>
    <item>
      <title>Re: extract data into OneLake using Fabric</title>
      <link>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4101018#M4990</link>
      <description>&lt;DIV&gt;&lt;DIV&gt;&lt;DIV class=""&gt;&lt;SPAN class=""&gt;How can I write to my Fabric Lakehouse via external app or Systems?&lt;/SPAN&gt;&lt;/DIV&gt;&lt;/DIV&gt;&lt;/DIV&gt;&lt;DIV class=""&gt;&lt;DIV class=""&gt;&lt;DIV class=""&gt;&lt;SPAN&gt;We are currently integrating multiple systems and need to push files directly into the Fabric Lakehouse. However, we would like to avoid using notebooks or pipelines for this process. Additionally, we anticipate adding more systems in the future and want to ensure the solution can scale accordingly.&lt;/SPAN&gt;&lt;/DIV&gt;&lt;DIV class=""&gt;&lt;SPAN&gt;Could you please advise on the best approach to achieve this?&lt;/SPAN&gt;&lt;/DIV&gt;&lt;/DIV&gt;&lt;/DIV&gt;</description>
      <pubDate>Wed, 14 Aug 2024 21:22:49 GMT</pubDate>
      <guid>https://community.fabric.microsoft.com/t5/Pipelines/extract-data-into-OneLake-using-Fabric/m-p/4101018#M4990</guid>
      <dc:creator>esraaE</dc:creator>
      <dc:date>2024-08-14T21:22:49Z</dc:date>
    </item>
  </channel>
</rss>

