<?xml version="1.0" encoding="UTF-8"?>
<rss xmlns:content="http://purl.org/rss/1.0/modules/content/" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:rdf="http://www.w3.org/1999/02/22-rdf-syntax-ns#" xmlns:taxo="http://purl.org/rss/1.0/modules/taxonomy/" version="2.0">
  <channel>
    <title>topic Re: How to subset / filter data based on frequency of unique observations within a variable in New SAS User</title>
    <link>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551953#M9088</link>
    <description>&lt;P&gt;SQL is very convenient in my opinion&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;PRE&gt;&lt;CODE class=" language-sas"&gt;
data have;
input  Obs        Number              Letter $;
cards;
1                36                      aaa
2                54                      aaa
3                34                      bbb
4                27                      aaa
5                58                      ccc
6                60                      aaa
7                39                      ddd
8                16                      bbb
9               12                       bbb
10              86                      eee
;

proc sql;
create table want as
select *
from have
group by letter
having n(letter)&amp;gt;=2
order by obs;
quit;&lt;/CODE&gt;&lt;/PRE&gt;</description>
    <pubDate>Thu, 18 Apr 2019 01:12:18 GMT</pubDate>
    <dc:creator>novinosrin</dc:creator>
    <dc:date>2019-04-18T01:12:18Z</dc:date>
    <item>
      <title>How to subset / filter data based on frequency of unique observations within a variable</title>
      <link>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551951#M9086</link>
      <description>&lt;P&gt;I am hoping to subset / filter a dataset by limiting based on the frequency of unique observations within a (character) variable although I'm not sure what procedure to use.&lt;/P&gt;&lt;P&gt;&amp;nbsp;&lt;/P&gt;&lt;P&gt;In the following example dataset:&lt;/P&gt;&lt;P&gt;&amp;nbsp;&lt;/P&gt;&lt;P&gt;Obs&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; Number&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; Letter&lt;/P&gt;&lt;P&gt;1&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 36&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; aaa&lt;/P&gt;&lt;P&gt;2&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 54&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; aaa&lt;/P&gt;&lt;P&gt;3&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 34&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; bbb&lt;/P&gt;&lt;P&gt;4&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 27&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; aaa&lt;/P&gt;&lt;P&gt;5&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 58&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; ccc&lt;/P&gt;&lt;P&gt;6&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 60&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; aaa&lt;/P&gt;&lt;P&gt;7&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 39&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; ddd&lt;/P&gt;&lt;P&gt;8&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 16&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; bbb&lt;/P&gt;&lt;P&gt;9&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp;12&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp;bbb&lt;/P&gt;&lt;P&gt;10&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; 86&amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; &amp;nbsp; eee&lt;/P&gt;&lt;P&gt;&amp;nbsp;&lt;/P&gt;&lt;P&gt;Is there a way to create a new dataset which includes the rows where the frequency of unique observation in 'Letter' is &amp;gt;/=2? (Ie - this would include rows with aaa (freq = 4) and bbb (freq = 3), which works out to be rows 1-4, 6, 8 &amp;amp; 9. I suppose I could use 'proc freq' to determine the frequency of each unique observation and manually 'drop' the values with &amp;lt;2, but in a large dataset this becomes cumbersome. Thanks in advance - I appreciate any advice &lt;span class="lia-unicode-emoji" title=":slightly_smiling_face:"&gt;🙂&lt;/span&gt;&lt;/P&gt;</description>
      <pubDate>Thu, 18 Apr 2019 00:58:59 GMT</pubDate>
      <guid>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551951#M9086</guid>
      <dc:creator>bretthouston</dc:creator>
      <dc:date>2019-04-18T00:58:59Z</dc:date>
    </item>
    <item>
      <title>Re: How to subset / filter data based on frequency of unique observations within a variable</title>
      <link>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551952#M9087</link>
      <description>&lt;P&gt;No to proc freq eh?&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;So what next?&lt;/P&gt;
&lt;P&gt;SQL?&lt;/P&gt;
&lt;P&gt;or Datastep?&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;</description>
      <pubDate>Thu, 18 Apr 2019 01:06:17 GMT</pubDate>
      <guid>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551952#M9087</guid>
      <dc:creator>novinosrin</dc:creator>
      <dc:date>2019-04-18T01:06:17Z</dc:date>
    </item>
    <item>
      <title>Re: How to subset / filter data based on frequency of unique observations within a variable</title>
      <link>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551953#M9088</link>
      <description>&lt;P&gt;SQL is very convenient in my opinion&lt;/P&gt;
&lt;P&gt;&amp;nbsp;&lt;/P&gt;
&lt;PRE&gt;&lt;CODE class=" language-sas"&gt;
data have;
input  Obs        Number              Letter $;
cards;
1                36                      aaa
2                54                      aaa
3                34                      bbb
4                27                      aaa
5                58                      ccc
6                60                      aaa
7                39                      ddd
8                16                      bbb
9               12                       bbb
10              86                      eee
;

proc sql;
create table want as
select *
from have
group by letter
having n(letter)&amp;gt;=2
order by obs;
quit;&lt;/CODE&gt;&lt;/PRE&gt;</description>
      <pubDate>Thu, 18 Apr 2019 01:12:18 GMT</pubDate>
      <guid>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551953#M9088</guid>
      <dc:creator>novinosrin</dc:creator>
      <dc:date>2019-04-18T01:12:18Z</dc:date>
    </item>
    <item>
      <title>Re: How to subset / filter data based on frequency of unique observations within a variable</title>
      <link>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551954#M9089</link>
      <description>&lt;P&gt;Thank-you! I just tried it and it worked perfectly. I really appreciate the help!&lt;/P&gt;</description>
      <pubDate>Thu, 18 Apr 2019 01:23:27 GMT</pubDate>
      <guid>https://communities.sas.com/t5/New-SAS-User/How-to-subset-filter-data-based-on-frequency-of-unique/m-p/551954#M9089</guid>
      <dc:creator>bretthouston</dc:creator>
      <dc:date>2019-04-18T01:23:27Z</dc:date>
    </item>
  </channel>
</rss>

