add offset param & email extraction (#51 )

* add offset param * [enh]: extract emails
chore: version number
2026-03-05 03:54:31 -08:00 · 2023-09-28 18:11:28 -05:00 · 2023-09-21 20:28:57 -05:00 · 2023-09-21 20:27:04 -05:00 · 2023-09-21 17:42:24 -05:00 · 2023-09-21 09:52:18 -07:00
24 changed files with 2085 additions and 2163 deletions
--- a/.github/workflows/publish-to-pypi.yml
+++ b/.github/workflows/publish-to-pypi.yml
@@ -7,27 +7,27 @@ jobs:
    runs-on: ubuntu-latest

    steps:
-    - uses: actions/checkout@v3
-    - name: Set up Python
-      uses: actions/setup-python@v4
-      with:
-        python-version: "3.10"
+      - uses: actions/checkout@v3
+      - name: Set up Python
+        uses: actions/setup-python@v4
+        with:
+          python-version: "3.10"

-    - name: Install poetry
-      run: >-
-        python3 -m
-        pip install
-        poetry
-        --user
+      - name: Install poetry
+        run: >-
+          python3 -m
+          pip install
+          poetry
+          --user

-    - name: Build distribution 📦
-      run: >-
-        python3 -m
-        poetry
-        build
+      - name: Build distribution 📦
+        run: >-
+          python3 -m
+          poetry
+          build

-    - name: Publish distribution 📦 to PyPI
-      if: startsWith(github.ref, 'refs/tags')
-      uses: pypa/gh-action-pypi-publish@release/v1
-      with:
-        password: ${{ secrets.PYPI_API_TOKEN }}
+      - name: Publish distribution 📦 to PyPI
+        if: startsWith(github.ref, 'refs/tags')
+        uses: pypa/gh-action-pypi-publish@release/v1
+        with:
+          password: ${{ secrets.PYPI_API_TOKEN }}
--- a/.gitignore
+++ b/.gitignore
@@ -1,9 +1,10 @@
-/.idea
-**/.DS_Store
 /venv/
-/ven/
+/.idea
 **/__pycache__/
+**/.pytest_cache/
+/.ipynb_checkpoints/
+**/output/
+**/.DS_Store
 *.pyc
 .env
-dist
-/.ipynb_checkpoints/
+dist
--- a/JobSpy_Demo.ipynb
+++ b/JobSpy_Demo.ipynb
@@ -1,701 +0,0 @@
-{
- "cells": [
-  {
-   "cell_type": "code",
-   "execution_count": 2,
-   "id": "c3f21577-477d-451e-9914-5d67e8a89075",
-   "metadata": {
-    "scrolled": true
-   },
-   "outputs": [
-    {
-     "data": {
-      "text/html": [
-       "<div>\n",
-       "<style scoped>\n",
-       "    .dataframe tbody tr th:only-of-type {\n",
-       "        vertical-align: middle;\n",
-       "    }\n",
-       "\n",
-       "    .dataframe tbody tr th {\n",
-       "        vertical-align: top;\n",
-       "    }\n",
-       "\n",
-       "    .dataframe thead th {\n",
-       "        text-align: right;\n",
-       "    }\n",
-       "</style>\n",
-       "<table border=\"1\" class=\"dataframe\">\n",
-       "  <thead>\n",
-       "    <tr style=\"text-align: right;\">\n",
-       "      <th></th>\n",
-       "      <th>site</th>\n",
-       "      <th>title</th>\n",
-       "      <th>company_name</th>\n",
-       "      <th>city</th>\n",
-       "      <th>state</th>\n",
-       "      <th>job_type</th>\n",
-       "      <th>interval</th>\n",
-       "      <th>min_amount</th>\n",
-       "      <th>max_amount</th>\n",
-       "      <th>job_url</th>\n",
-       "      <th>description</th>\n",
-       "    </tr>\n",
-       "  </thead>\n",
-       "  <tbody>\n",
-       "    <tr>\n",
-       "      <th>0</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Firmware Engineer</td>\n",
-       "      <td>Advanced Motion Controls</td>\n",
-       "      <td>Camarillo</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>145000</td>\n",
-       "      <td>110000</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=a2e7077fdd3c...</td>\n",
-       "      <td>We are looking for an experienced Firmware Eng...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>1</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Computer Engineer</td>\n",
-       "      <td>Honeywell</td>\n",
-       "      <td></td>\n",
-       "      <td>None</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=5a1da623ee75...</td>\n",
-       "      <td>Join a team recognized for leadership, innovat...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>2</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Splunk</td>\n",
-       "      <td>Remote</td>\n",
-       "      <td>None</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>159500</td>\n",
-       "      <td>116000</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=155495ca3f46...</td>\n",
-       "      <td>A little about us. Splunk is the key to enterp...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>3</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Development Operations Engineer</td>\n",
-       "      <td>Stratacache</td>\n",
-       "      <td>Dayton</td>\n",
-       "      <td>OH</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>90000</td>\n",
-       "      <td>83573</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=77cf3540c06e...</td>\n",
-       "      <td>Stratacache, Inc. delivers in-store retail exp...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>4</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Computer Engineer</td>\n",
-       "      <td>Honeywell</td>\n",
-       "      <td></td>\n",
-       "      <td>None</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=7fadbb7c936f...</td>\n",
-       "      <td>Join a team recognized for leadership, innovat...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>5</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Full Stack Developer</td>\n",
-       "      <td>Reinventing Geospatial, Inc. (RGi)</td>\n",
-       "      <td>Herndon</td>\n",
-       "      <td>VA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=11b2b5b0dd44...</td>\n",
-       "      <td>Job Highlights As a Full Stack Software Engine...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>6</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Workiva</td>\n",
-       "      <td>Remote</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>134000</td>\n",
-       "      <td>79000</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=ec3ab6eb9253...</td>\n",
-       "      <td>Are you ready to embark on an exciting journey...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>7</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Senior Software Engineer</td>\n",
-       "      <td>SciTec</td>\n",
-       "      <td>Boulder</td>\n",
-       "      <td>CO</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>164000</td>\n",
-       "      <td>93000</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=781e4cf0cf6d...</td>\n",
-       "      <td>SciTec has been awarded multiple government co...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>8</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Microsoft</td>\n",
-       "      <td></td>\n",
-       "      <td>None</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>182600</td>\n",
-       "      <td>94300</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=21e05b9e9d96...</td>\n",
-       "      <td>At Microsoft we are seeking people who have a ...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>9</th>\n",
-       "      <td>indeed</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Avalon Healthcare Solutions</td>\n",
-       "      <td>Remote</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.indeed.com/viewjob?jk=da35b9bb74a0...</td>\n",
-       "      <td>Avalon Healthcare Solutions, headquartered in ...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>10</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Fieldguide</td>\n",
-       "      <td>San Francisco</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3696158160</td>\n",
-       "      <td>About us:Fieldguide is establishing a new stat...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>11</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Sunnyvale</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3693012711</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>12</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Edwards</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3700669785</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>13</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Fort Worth</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3701775201</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>14</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Fort Worth</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3701772329</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>15</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Fort Worth</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3701769637</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>16</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>SpiderOak</td>\n",
-       "      <td>Austin</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3707174719</td>\n",
-       "      <td>We're only as strong as our weakest link.In th...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>17</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer - Early Career</td>\n",
-       "      <td>Lockheed Martin</td>\n",
-       "      <td>Fort Worth</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3701770659</td>\n",
-       "      <td>Description:By bringing together people that u...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>18</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Full-Stack Software Engineer</td>\n",
-       "      <td>Rain</td>\n",
-       "      <td>New York</td>\n",
-       "      <td>NY</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3696158877</td>\n",
-       "      <td>Rain’s mission is to create the fastest and ea...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>19</th>\n",
-       "      <td>linkedin</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Nike</td>\n",
-       "      <td>Portland</td>\n",
-       "      <td>OR</td>\n",
-       "      <td>contract</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://www.linkedin.com/jobs/view/3693340247</td>\n",
-       "      <td>Work options: FlexibleWe consider remote, on-p...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>20</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>(USA) Software Engineer III - Prototype Engine...</td>\n",
-       "      <td>Walmart</td>\n",
-       "      <td>Dallas</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>None</td>\n",
-       "      <td>https://click.appcast.io/track/hcgsw4k?cs=ngp&amp;...</td>\n",
-       "      <td>We are currently seeking a highly skilled and ...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>21</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Engineer - New Grad</td>\n",
-       "      <td>ZipRecruiter</td>\n",
-       "      <td>Santa Monica</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>130000</td>\n",
-       "      <td>150000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/ziprecruiter...</td>\n",
-       "      <td>We offer a hybrid work environment. Most US-ba...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>22</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Developer</td>\n",
-       "      <td>Robert Half</td>\n",
-       "      <td>Corpus Christi</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>105000</td>\n",
-       "      <td>115000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/robert-half-...</td>\n",
-       "      <td>Robert Half has an opening for a Software Deve...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>23</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Engineer</td>\n",
-       "      <td>Advantage Technical</td>\n",
-       "      <td>Ontario</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>100000</td>\n",
-       "      <td>150000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/advantage-te...</td>\n",
-       "      <td>New career opportunity available with major Ma...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>24</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Developer</td>\n",
-       "      <td>Robert Half</td>\n",
-       "      <td>Tucson</td>\n",
-       "      <td>AZ</td>\n",
-       "      <td>temporary</td>\n",
-       "      <td>hourly</td>\n",
-       "      <td>47</td>\n",
-       "      <td>55</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/robert-half-...</td>\n",
-       "      <td>Robert Half is accepting inquiries for a SQL S...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>25</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Full Stack Software Engineer</td>\n",
-       "      <td>ZipRecruiter</td>\n",
-       "      <td>Phoenix</td>\n",
-       "      <td>AZ</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>105000</td>\n",
-       "      <td>145000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/ziprecruiter...</td>\n",
-       "      <td>We offer a hybrid work environment. Most US-ba...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>26</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Developer IV</td>\n",
-       "      <td>Kforce Inc.</td>\n",
-       "      <td>Mountain View</td>\n",
-       "      <td>CA</td>\n",
-       "      <td>contract</td>\n",
-       "      <td>hourly</td>\n",
-       "      <td>55</td>\n",
-       "      <td>75</td>\n",
-       "      <td>https://www.kforce.com/Jobs/job.aspx?job=1696~...</td>\n",
-       "      <td>Kforce has a client that is seeking a Software...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>27</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Developer | Onsite | Omaha, NE - Omaha</td>\n",
-       "      <td>OneStaff Medical</td>\n",
-       "      <td>Omaha</td>\n",
-       "      <td>NE</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>60000</td>\n",
-       "      <td>110000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/onestaff-med...</td>\n",
-       "      <td>Company Description: We are looking for a well...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>28</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Senior Software Engineer</td>\n",
-       "      <td>RightStaff, Inc.</td>\n",
-       "      <td>Dallas</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>120000</td>\n",
-       "      <td>180000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/rightstaff-i...</td>\n",
-       "      <td>Job Description:We are seeking a talented and ...</td>\n",
-       "    </tr>\n",
-       "    <tr>\n",
-       "      <th>29</th>\n",
-       "      <td>zip_recruiter</td>\n",
-       "      <td>Software Developer - .Net Core - 12886</td>\n",
-       "      <td>Walker Elliott</td>\n",
-       "      <td>Dallas</td>\n",
-       "      <td>TX</td>\n",
-       "      <td>fulltime</td>\n",
-       "      <td>yearly</td>\n",
-       "      <td>105000</td>\n",
-       "      <td>130000</td>\n",
-       "      <td>https://www.ziprecruiter.com/jobs/walker-ellio...</td>\n",
-       "      <td>Our highly successful DFW based client has bee...</td>\n",
-       "    </tr>\n",
-       "  </tbody>\n",
-       "</table>\n",
-       "</div>"
-      ],
-      "text/plain": [
-       "             site                                              title  \\\n",
-       "0          indeed                                  Firmware Engineer   \n",
-       "1          indeed                                  Computer Engineer   \n",
-       "2          indeed                                  Software Engineer   \n",
-       "3          indeed                    Development Operations Engineer   \n",
-       "4          indeed                                  Computer Engineer   \n",
-       "5          indeed                               Full Stack Developer   \n",
-       "6          indeed                                  Software Engineer   \n",
-       "7          indeed                           Senior Software Engineer   \n",
-       "8          indeed                                  Software Engineer   \n",
-       "9          indeed                                  Software Engineer   \n",
-       "10       linkedin                                  Software Engineer   \n",
-       "11       linkedin                   Software Engineer - Early Career   \n",
-       "12       linkedin                   Software Engineer - Early Career   \n",
-       "13       linkedin                   Software Engineer - Early Career   \n",
-       "14       linkedin                   Software Engineer - Early Career   \n",
-       "15       linkedin                   Software Engineer - Early Career   \n",
-       "16       linkedin                                  Software Engineer   \n",
-       "17       linkedin                   Software Engineer - Early Career   \n",
-       "18       linkedin                       Full-Stack Software Engineer   \n",
-       "19       linkedin                                  Software Engineer   \n",
-       "20  zip_recruiter  (USA) Software Engineer III - Prototype Engine...   \n",
-       "21  zip_recruiter                       Software Engineer - New Grad   \n",
-       "22  zip_recruiter                                 Software Developer   \n",
-       "23  zip_recruiter                                  Software Engineer   \n",
-       "24  zip_recruiter                                 Software Developer   \n",
-       "25  zip_recruiter                       Full Stack Software Engineer   \n",
-       "26  zip_recruiter                              Software Developer IV   \n",
-       "27  zip_recruiter    Software Developer | Onsite | Omaha, NE - Omaha   \n",
-       "28  zip_recruiter                           Senior Software Engineer   \n",
-       "29  zip_recruiter             Software Developer - .Net Core - 12886   \n",
-       "\n",
-       "                          company_name            city state   job_type  \\\n",
-       "0             Advanced Motion Controls       Camarillo    CA   fulltime   \n",
-       "1                            Honeywell                  None   fulltime   \n",
-       "2                               Splunk          Remote  None   fulltime   \n",
-       "3                          Stratacache          Dayton    OH   fulltime   \n",
-       "4                            Honeywell                  None   fulltime   \n",
-       "5   Reinventing Geospatial, Inc. (RGi)         Herndon    VA   fulltime   \n",
-       "6                              Workiva          Remote  None       None   \n",
-       "7                               SciTec         Boulder    CO   fulltime   \n",
-       "8                            Microsoft                  None   fulltime   \n",
-       "9          Avalon Healthcare Solutions          Remote  None       None   \n",
-       "10                          Fieldguide   San Francisco    CA   fulltime   \n",
-       "11                     Lockheed Martin       Sunnyvale    CA   fulltime   \n",
-       "12                     Lockheed Martin         Edwards    CA   fulltime   \n",
-       "13                     Lockheed Martin      Fort Worth    TX   fulltime   \n",
-       "14                     Lockheed Martin      Fort Worth    TX   fulltime   \n",
-       "15                     Lockheed Martin      Fort Worth    TX   fulltime   \n",
-       "16                           SpiderOak          Austin    TX   fulltime   \n",
-       "17                     Lockheed Martin      Fort Worth    TX   fulltime   \n",
-       "18                                Rain        New York    NY   fulltime   \n",
-       "19                                Nike        Portland    OR   contract   \n",
-       "20                             Walmart          Dallas    TX       None   \n",
-       "21                        ZipRecruiter    Santa Monica    CA   fulltime   \n",
-       "22                         Robert Half  Corpus Christi    TX   fulltime   \n",
-       "23                 Advantage Technical         Ontario    CA   fulltime   \n",
-       "24                         Robert Half          Tucson    AZ  temporary   \n",
-       "25                        ZipRecruiter         Phoenix    AZ   fulltime   \n",
-       "26                         Kforce Inc.   Mountain View    CA   contract   \n",
-       "27                    OneStaff Medical           Omaha    NE   fulltime   \n",
-       "28                    RightStaff, Inc.          Dallas    TX   fulltime   \n",
-       "29                      Walker Elliott          Dallas    TX   fulltime   \n",
-       "\n",
-       "   interval min_amount max_amount  \\\n",
-       "0    yearly     145000     110000   \n",
-       "1      None       None       None   \n",
-       "2    yearly     159500     116000   \n",
-       "3    yearly      90000      83573   \n",
-       "4      None       None       None   \n",
-       "5      None       None       None   \n",
-       "6    yearly     134000      79000   \n",
-       "7    yearly     164000      93000   \n",
-       "8    yearly     182600      94300   \n",
-       "9      None       None       None   \n",
-       "10   yearly       None       None   \n",
-       "11   yearly       None       None   \n",
-       "12   yearly       None       None   \n",
-       "13   yearly       None       None   \n",
-       "14   yearly       None       None   \n",
-       "15   yearly       None       None   \n",
-       "16   yearly       None       None   \n",
-       "17   yearly       None       None   \n",
-       "18   yearly       None       None   \n",
-       "19   yearly       None       None   \n",
-       "20     None       None       None   \n",
-       "21   yearly     130000     150000   \n",
-       "22   yearly     105000     115000   \n",
-       "23   yearly     100000     150000   \n",
-       "24   hourly         47         55   \n",
-       "25   yearly     105000     145000   \n",
-       "26   hourly         55         75   \n",
-       "27   yearly      60000     110000   \n",
-       "28   yearly     120000     180000   \n",
-       "29   yearly     105000     130000   \n",
-       "\n",
-       "                                              job_url  \\\n",
-       "0   https://www.indeed.com/viewjob?jk=a2e7077fdd3c...   \n",
-       "1   https://www.indeed.com/viewjob?jk=5a1da623ee75...   \n",
-       "2   https://www.indeed.com/viewjob?jk=155495ca3f46...   \n",
-       "3   https://www.indeed.com/viewjob?jk=77cf3540c06e...   \n",
-       "4   https://www.indeed.com/viewjob?jk=7fadbb7c936f...   \n",
-       "5   https://www.indeed.com/viewjob?jk=11b2b5b0dd44...   \n",
-       "6   https://www.indeed.com/viewjob?jk=ec3ab6eb9253...   \n",
-       "7   https://www.indeed.com/viewjob?jk=781e4cf0cf6d...   \n",
-       "8   https://www.indeed.com/viewjob?jk=21e05b9e9d96...   \n",
-       "9   https://www.indeed.com/viewjob?jk=da35b9bb74a0...   \n",
-       "10      https://www.linkedin.com/jobs/view/3696158160   \n",
-       "11      https://www.linkedin.com/jobs/view/3693012711   \n",
-       "12      https://www.linkedin.com/jobs/view/3700669785   \n",
-       "13      https://www.linkedin.com/jobs/view/3701775201   \n",
-       "14      https://www.linkedin.com/jobs/view/3701772329   \n",
-       "15      https://www.linkedin.com/jobs/view/3701769637   \n",
-       "16      https://www.linkedin.com/jobs/view/3707174719   \n",
-       "17      https://www.linkedin.com/jobs/view/3701770659   \n",
-       "18      https://www.linkedin.com/jobs/view/3696158877   \n",
-       "19      https://www.linkedin.com/jobs/view/3693340247   \n",
-       "20  https://click.appcast.io/track/hcgsw4k?cs=ngp&...   \n",
-       "21  https://www.ziprecruiter.com/jobs/ziprecruiter...   \n",
-       "22  https://www.ziprecruiter.com/jobs/robert-half-...   \n",
-       "23  https://www.ziprecruiter.com/jobs/advantage-te...   \n",
-       "24  https://www.ziprecruiter.com/jobs/robert-half-...   \n",
-       "25  https://www.ziprecruiter.com/jobs/ziprecruiter...   \n",
-       "26  https://www.kforce.com/Jobs/job.aspx?job=1696~...   \n",
-       "27  https://www.ziprecruiter.com/jobs/onestaff-med...   \n",
-       "28  https://www.ziprecruiter.com/jobs/rightstaff-i...   \n",
-       "29  https://www.ziprecruiter.com/jobs/walker-ellio...   \n",
-       "\n",
-       "                                          description  \n",
-       "0   We are looking for an experienced Firmware Eng...  \n",
-       "1   Join a team recognized for leadership, innovat...  \n",
-       "2   A little about us. Splunk is the key to enterp...  \n",
-       "3   Stratacache, Inc. delivers in-store retail exp...  \n",
-       "4   Join a team recognized for leadership, innovat...  \n",
-       "5   Job Highlights As a Full Stack Software Engine...  \n",
-       "6   Are you ready to embark on an exciting journey...  \n",
-       "7   SciTec has been awarded multiple government co...  \n",
-       "8   At Microsoft we are seeking people who have a ...  \n",
-       "9   Avalon Healthcare Solutions, headquartered in ...  \n",
-       "10  About us:Fieldguide is establishing a new stat...  \n",
-       "11  Description:By bringing together people that u...  \n",
-       "12  Description:By bringing together people that u...  \n",
-       "13  Description:By bringing together people that u...  \n",
-       "14  Description:By bringing together people that u...  \n",
-       "15  Description:By bringing together people that u...  \n",
-       "16  We're only as strong as our weakest link.In th...  \n",
-       "17  Description:By bringing together people that u...  \n",
-       "18  Rain’s mission is to create the fastest and ea...  \n",
-       "19  Work options: FlexibleWe consider remote, on-p...  \n",
-       "20  We are currently seeking a highly skilled and ...  \n",
-       "21  We offer a hybrid work environment. Most US-ba...  \n",
-       "22  Robert Half has an opening for a Software Deve...  \n",
-       "23  New career opportunity available with major Ma...  \n",
-       "24  Robert Half is accepting inquiries for a SQL S...  \n",
-       "25  We offer a hybrid work environment. Most US-ba...  \n",
-       "26  Kforce has a client that is seeking a Software...  \n",
-       "27  Company Description: We are looking for a well...  \n",
-       "28  Job Description:We are seeking a talented and ...  \n",
-       "29  Our highly successful DFW based client has bee...  "
-      ]
-     },
-     "metadata": {},
-     "output_type": "display_data"
-    }
-   ],
-   "source": [
-    "from jobspy import scrape_jobs\n",
-    "import pandas as pd\n",
-    "\n",
-    "jobs: pd.DataFrame = scrape_jobs(\n",
-    "    site_name=[\"indeed\", \"linkedin\", \"zip_recruiter\"],\n",
-    "    search_term=\"software engineer\",\n",
-    "    results_wanted=10\n",
-    ")\n",
-    "\n",
-    "if jobs.empty:\n",
-    "    print(\"No jobs found.\")\n",
-    "else:\n",
-    "    # 1 print\n",
-    "    pd.set_option('display.max_columns', None)\n",
-    "    pd.set_option('display.max_rows', None)\n",
-    "    pd.set_option('display.width', None)\n",
-    "    pd.set_option('display.max_colwidth', 50)  # set to 0 to see full job url / desc\n",
-    "    print(jobs)\n",
-    "\n",
-    "    # 2 display in Jupyter Notebook\n",
-    "    display(jobs)\n",
-    "\n",
-    "    # 3 output to csv\n",
-    "    jobs.to_csv('jobs.csv', index=False)"
-   ]
-  },
-  {
-   "cell_type": "code",
-   "execution_count": null,
-   "id": "efd667ef-fdf0-452a-b5e5-ce6825755be7",
-   "metadata": {},
-   "outputs": [],
-   "source": []
-  },
-  {
-   "cell_type": "code",
-   "execution_count": null,
-   "id": "1574dc17-0a42-4655-964f-5c03a6d3deb0",
-   "metadata": {},
-   "outputs": [],
-   "source": []
-  }
- ],
- "metadata": {
-  "kernelspec": {
-   "display_name": "my-poetry-env",
-   "language": "python",
-   "name": "my-poetry-env"
-  },
-  "language_info": {
-   "codemirror_mode": {
-    "name": "ipython",
-    "version": 3
-   },
-   "file_extension": ".py",
-   "mimetype": "text/x-python",
-   "name": "python",
-   "nbconvert_exporter": "python",
-   "pygments_lexer": "ipython3",
-   "version": "3.10.11"
-  }
- },
- "nbformat": 4,
- "nbformat_minor": 5
-}
--- a/README.md
+++ b/README.md
@@ -1,15 +1,33 @@
-# JobSpy
+<img src="https://github.com/cullenwatson/JobSpy/assets/78247585/ae185b7e-e444-4712-8bb9-fa97f53e896b" width="400">

 **JobSpy** is a simple, yet comprehensive, job scraping library.
+
+**Not technical?** Try out the web scraping tool on our site at [usejobspy.com](https://usejobspy.com).
+
+*Looking to build a data-focused software product?* **[Book a call](https://calendly.com/zachary-products/15min)** *to
+work with us.*  
+\
+Check out another project we wrote: ***[HomeHarvest](https://github.com/ZacharyHampton/HomeHarvest)** – a Python package
+for real estate scraping*
+
 ## Features

 - Scrapes job postings from **LinkedIn**, **Indeed** & **ZipRecruiter** simultaneously
 - Aggregates the job postings in a Pandas DataFrame
+- Proxy support (HTTP/S, SOCKS)
+
+[Video Guide for JobSpy](https://www.youtube.com/watch?v=RuP1HrAZnxs&pp=ygUgam9icyBzY3JhcGVyIGJvdCBsaW5rZWRpbiBpbmRlZWQ%3D) -
+Updated for release v1.1.3
+
+![jobspy](https://github.com/cullenwatson/JobSpy/assets/78247585/ec7ef355-05f6-4fd3-8161-a817e31c5c57)

 ### Installation
-`pip install python-jobspy`  
-  
-  _Python version >= [3.10](https://www.python.org/downloads/release/python-3100/) required_ 
+
+```
+pip install --upgrade python-jobspy
+```
+
+_Python version >= [3.10](https://www.python.org/downloads/release/python-3100/) required_

 ### Usage

@@ -20,38 +38,49 @@ import pandas as pd
 jobs: pd.DataFrame = scrape_jobs(
    site_name=["indeed", "linkedin", "zip_recruiter"],
    search_term="software engineer",
-    results_wanted=10
+    location="Dallas, TX",
+    results_wanted=10,
+
+    country_indeed='USA'  # only needed for indeed
+
+    # use if you want to use a proxy
+    # proxy="http://jobspy:5a4vpWtj8EeJ2hoYzk@ca.smartproxy.com:20001",
+    # offset=25 # use if you want to start at a specific offset
 )

-if jobs.empty:
-    print("No jobs found.")
-else:
-    # 1 print
-    pd.set_option('display.max_columns', None)
-    pd.set_option('display.max_rows', None)
-    pd.set_option('display.width', None)
-    pd.set_option('display.max_colwidth', 50)  # set to 0 to see full job url / desc
-    print(jobs)
+# formatting for pandas
+pd.set_option('display.max_columns', None)
+pd.set_option('display.max_rows', None)
+pd.set_option('display.width', None)
+pd.set_option('display.max_colwidth', 50)  # set to 0 to see full job url / desc

-    # 2 display in Jupyter Notebook
-    # display(jobs)
+# 1 output to console
+print(jobs)
+
+# 2 display in Jupyter Notebook (1. pip install jupyter 2. jupyter notebook)
+# display(jobs)
+
+# 3 output to .csv
+# jobs.to_csv('jobs.csv', index=False)
+
+# 4 output to .xlsx
+# jobs.to_xlsx('jobs.xlsx', index=False)

-    # 3 output to csv
-    # jobs.to_csv('jobs.csv', index=False)
 ```

 ### Output
-```
-             site                                              title                    company_name                 city state   job_type interval min_amount max_amount                                            job_url                                        description
-           indeed                                  Software Engineer                AMERICAN SYSTEMS            Arlington    VA       None   yearly     200000     150000  https://www.indeed.com/viewjob?jk=5e409e577046...  THIS POSITION COMES WITH A 10K SIGNING BONUS! ...
-           indeed                           Senior Software Engineer                TherapyNotes.com         Philadelphia    PA   fulltime   yearly     135000     110000  https://www.indeed.com/viewjob?jk=da39574a40cb...  About Us TherapyNotes is the national leader i...
-         linkedin                   Software Engineer - Early Career                 Lockheed Martin            Sunnyvale    CA   fulltime   yearly       None       None      https://www.linkedin.com/jobs/view/3693012711  Description:By bringing together people that u...
-         linkedin                       Full-Stack Software Engineer                            Rain             New York    NY   fulltime   yearly       None       None      https://www.linkedin.com/jobs/view/3696158877  Rain’s mission is to create the fastest and ea...
-    zip_recruiter                       Software Engineer - New Grad                    ZipRecruiter         Santa Monica    CA   fulltime   yearly     130000     150000  https://www.ziprecruiter.com/jobs/ziprecruiter...  We offer a hybrid work environment. Most US-ba...
-    zip_recruiter                                 Software Developer                      TEKsystems              Phoenix    AZ   fulltime   hourly         65         75  https://www.ziprecruiter.com/jobs/teksystems-0...  Top Skills' Details• 6 years of Java developme.```
-```
-### Parameters for `scrape_jobs()`

+```
+SITE           TITLE                             COMPANY_NAME      CITY          STATE  JOB_TYPE  INTERVAL  MIN_AMOUNT  MAX_AMOUNT  JOB_URL                                            DESCRIPTION
+indeed         Software Engineer                 AMERICAN SYSTEMS  Arlington     VA     None      yearly    200000      150000      https://www.indeed.com/viewjob?jk=5e409e577046...  THIS POSITION COMES WITH A 10K SIGNING BONUS!...
+indeed         Senior Software Engineer          TherapyNotes.com  Philadelphia  PA     fulltime  yearly    135000      110000      https://www.indeed.com/viewjob?jk=da39574a40cb...  About Us TherapyNotes is the national leader i...
+linkedin       Software Engineer - Early Career  Lockheed Martin   Sunnyvale     CA     fulltime  yearly    None        None        https://www.linkedin.com/jobs/view/3693012711      Description:By bringing together people that u...
+linkedin       Full-Stack Software Engineer      Rain              New York      NY     fulltime  yearly    None        None        https://www.linkedin.com/jobs/view/3696158877      Rain’s mission is to create the fastest and ea...
+zip_recruiter Software Engineer - New Grad       ZipRecruiter      Santa Monica  CA     fulltime  yearly    130000      150000      https://www.ziprecruiter.com/jobs/ziprecruiter...  We offer a hybrid work environment. Most US-ba...
+zip_recruiter Software Developer                 TEKsystems        Phoenix       AZ     fulltime  hourly    65          75          https://www.ziprecruiter.com/jobs/teksystems-0...  Top Skills' Details• 6 years of Java developme...
+```
+
+### Parameters for `scrape_jobs()`

 ```plaintext
 Required
@@ -61,39 +90,103 @@ Optional
 ├── location (int)
 ├── distance (int): in miles
 ├── job_type (enum): fulltime, parttime, internship, contract
+├── proxy (str): in format 'http://user:pass@host:port' or [https, socks]
 ├── is_remote (bool)
 ├── results_wanted (int): number of job results to retrieve for each site specified in 'site_type'
-├── easy_apply (bool): filters for jobs on LinkedIn that have the 'Easy Apply' option
+├── easy_apply (bool): filters for jobs that are hosted on LinkedIn
+├── country_indeed (enum): filters the country on Indeed (see below for correct spelling)
+├── offset (enum): starts the search from an offset (e.g. 25 will start the search from the 25th result)
 ```

-### Response Schema
+### JobPost Schema
+
 ```plaintext
 JobPost
 ├── title (str)
-├── company_name (str)
+├── company (str)
 ├── job_url (str)
 ├── location (object)
 │   ├── country (str)
 │   ├── city (str)
 │   ├── state (str)
 ├── description (str)
-├── job_type (enum)
+├── job_type (enum): fulltime, parttime, internship, contract
 ├── compensation (object)
-│   ├── interval (CompensationInterval): yearly, monthly, weekly, daily, hourly
-│   ├── min_amount (float)
-│   ├── max_amount (float)
-│   └── currency (str)
-└── date_posted (datetime)
-
+│   ├── interval (enum): yearly, monthly, weekly, daily, hourly
+│   ├── min_amount (int)
+│   ├── max_amount (int)
+│   └── currency (enum)
+└── date_posted (date)
 ```

+### Exceptions
+
+The following exceptions may be raised when using JobSpy:
+
+* `LinkedInException`
+* `IndeedException`
+* `ZipRecruiterException`
+
+## Supported Countries for Job Searching
+
+### **LinkedIn**
+
+LinkedIn searches globally & uses only the `location` parameter.
+
+### **ZipRecruiter**
+
+ZipRecruiter searches for jobs in **US/Canada** & uses only the `location` parameter.
+
+### **Indeed**
+
+Indeed supports most countries, but the `country_indeed` parameter is required. Additionally, use the `location`
+parameter to narrow down the location, e.g. city & state if necessary.
+
+You can specify the following countries when searching on Indeed (use the exact name):
+
+|                      |              |            |                |
+|----------------------|--------------|------------|----------------|
+| Argentina            | Australia    | Austria    | Bahrain        |
+| Belgium              | Brazil       | Canada     | Chile          |
+| China                | Colombia     | Costa Rica | Czech Republic |
+| Denmark              | Ecuador      | Egypt      | Finland        |
+| France               | Germany      | Greece     | Hong Kong      |
+| Hungary              | India        | Indonesia  | Ireland        |
+| Israel               | Italy        | Japan      | Kuwait         |
+| Luxembourg           | Malaysia     | Mexico     | Morocco        |
+| Netherlands          | New Zealand  | Nigeria    | Norway         |
+| Oman                 | Pakistan     | Panama     | Peru           |
+| Philippines          | Poland       | Portugal   | Qatar          |
+| Romania              | Saudi Arabia | Singapore  | South Africa   |
+| South Korea          | Spain        | Sweden     | Switzerland    |
+| Taiwan               | Thailand     | Turkey     | Ukraine        |
+| United Arab Emirates | UK           | USA        | Uruguay        |
+| Venezuela            | Vietnam      |            |                |
+
+## Frequently Asked Questions
+
+---
+
+**Q: Encountering issues with your queries?**  
+**A:** Try reducing the number of `results_wanted` and/or broadening the filters. If problems
+persist, [submit an issue](https://github.com/cullenwatson/JobSpy/issues).
+
+---
+
+**Q: Received a response code 429?**  
+**A:** This indicates that you have been blocked by the job board site for sending too many requests. Currently, *
+*LinkedIn** is particularly aggressive with blocking. We recommend:
+
+- Waiting a few seconds between requests.
+- Trying a VPN or proxy to change your IP address.
+
+---
+
+**Q: Experiencing a "Segmentation fault: 11" on macOS Catalina?**  
+**A:** This is due to `tls_client` dependency not supporting your architecture. Solutions and workarounds include:
+
+- Upgrade to a newer version of MacOS
+- Reach out to the maintainers of [tls_client](https://github.com/bogdanfinn/tls-client) for fixes
+

-### FAQ
-  
-#### Encountering issues with your queries?
-  
-Try reducing the number of `results_wanted` and/or broadening the filters. If problems persist, please submit an issue.
-  
-#### Received a response code 429?
-This means you've been blocked by the job board site for sending too many requests. Consider waiting a few seconds, or try using a VPN. Proxy support coming soon.
  
--- a/examples/JobSpy_Demo.ipynb
+++ b/examples/JobSpy_Demo.ipynb
@@ -0,0 +1,167 @@
+{
+ "cells": [
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "00a94b47-f47b-420f-ba7e-714ef219c006",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "from jobspy import scrape_jobs\n",
+    "import pandas as pd\n",
+    "from IPython.display import display, HTML"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "9f773e6c-d9fc-42cc-b0ef-63b739e78435",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "pd.set_option('display.max_columns', None)\n",
+    "pd.set_option('display.max_rows', None)\n",
+    "pd.set_option('display.width', None)\n",
+    "pd.set_option('display.max_colwidth', 50)"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "1253c1f8-9437-492e-9dd3-e7fe51099420",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# example 1 (no hyperlinks, USA)\n",
+    "jobs = scrape_jobs(\n",
+    "    site_name=[\"linkedin\"],\n",
+    "    location='san francisco',\n",
+    "    search_term=\"engineer\",\n",
+    "    results_wanted=5,\n",
+    "\n",
+    "    # use if you want to use a proxy\n",
+    "    # proxy=\"socks5://jobspy:5a4vpWtj4EeJ2hoYzk@us.smartproxy.com:10001\",\n",
+    "    proxy=\"http://jobspy:5a4vpWtj4EeJ2hoYzk@us.smartproxy.com:10001\",\n",
+    "    #proxy=\"https://jobspy:5a4vpWtj4EeJ2hoYzk@us.smartproxy.com:10001\",\n",
+    ")\n",
+    "display(jobs)"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "6a581b2d-f7da-4fac-868d-9efe143ee20a",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# example 2 - remote USA & hyperlinks\n",
+    "jobs = scrape_jobs(\n",
+    "    site_name=[\"linkedin\", \"zip_recruiter\", \"indeed\"],\n",
+    "    # location='san francisco',\n",
+    "    search_term=\"software engineer\",\n",
+    "    country_indeed=\"USA\",\n",
+    "    hyperlinks=True,\n",
+    "    is_remote=True,\n",
+    "    results_wanted=5, \n",
+    ")"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "fe8289bc-5b64-4202-9a64-7c117c83fd9a",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# use if hyperlinks=True\n",
+    "html = jobs.to_html(escape=False)\n",
+    "# change max-width: 200px to show more or less of the content\n",
+    "truncate_width = f'<style>.dataframe td {{ max-width: 200px; overflow: hidden; text-overflow: ellipsis; white-space: nowrap; }}</style>{html}'\n",
+    "display(HTML(truncate_width))"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "951c2fe1-52ff-407d-8bb1-068049b36777",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# example 3 - with hyperlinks, international - linkedin (no zip_recruiter)\n",
+    "jobs = scrape_jobs(\n",
+    "    site_name=[\"linkedin\"],\n",
+    "    location='berlin',\n",
+    "    search_term=\"engineer\",\n",
+    "    hyperlinks=True,\n",
+    "    results_wanted=5,\n",
+    "    easy_apply=True\n",
+    ")"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "1e37a521-caef-441c-8fc2-2eb5b2e7da62",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# use if hyperlinks=True\n",
+    "html = jobs.to_html(escape=False)\n",
+    "# change max-width: 200px to show more or less of the content\n",
+    "truncate_width = f'<style>.dataframe td {{ max-width: 200px; overflow: hidden; text-overflow: ellipsis; white-space: nowrap; }}</style>{html}'\n",
+    "display(HTML(truncate_width))"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "0650e608-0b58-4bf5-ae86-68348035b16a",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# example 4 - international indeed (no zip_recruiter)\n",
+    "jobs = scrape_jobs(\n",
+    "    site_name=[\"indeed\"],\n",
+    "    search_term=\"engineer\",\n",
+    "    country_indeed = \"China\",\n",
+    "    hyperlinks=True\n",
+    ")"
+   ]
+  },
+  {
+   "cell_type": "code",
+   "execution_count": null,
+   "id": "40913ac8-3f8a-4d7e-ac47-afb88316432b",
+   "metadata": {},
+   "outputs": [],
+   "source": [
+    "# use if hyperlinks=True\n",
+    "html = jobs.to_html(escape=False)\n",
+    "# change max-width: 200px to show more or less of the content\n",
+    "truncate_width = f'<style>.dataframe td {{ max-width: 200px; overflow: hidden; text-overflow: ellipsis; white-space: nowrap; }}</style>{html}'\n",
+    "display(HTML(truncate_width))"
+   ]
+  }
+ ],
+ "metadata": {
+  "kernelspec": {
+   "display_name": "Python 3 (ipykernel)",
+   "language": "python",
+   "name": "python3"
+  },
+  "language_info": {
+   "codemirror_mode": {
+    "name": "ipython",
+    "version": 3
+   },
+   "file_extension": ".py",
+   "mimetype": "text/x-python",
+   "name": "python",
+   "nbconvert_exporter": "python",
+   "pygments_lexer": "ipython3",
+   "version": "3.11.5"
+  }
+ },
+ "nbformat": 4,
+ "nbformat_minor": 5
+}
--- a/examples/JobSpy_Demo.py
+++ b/examples/JobSpy_Demo.py
@@ -0,0 +1,31 @@
+from jobspy import scrape_jobs
+import pandas as pd
+
+jobs: pd.DataFrame = scrape_jobs(
+    site_name=["indeed", "linkedin", "zip_recruiter"],
+    search_term="software engineer",
+    location="Dallas, TX",
+    results_wanted=50,  # be wary the higher it is, the more likey you'll get blocked (rotating proxy should work tho)
+    country_indeed='USA',
+    offset=25  # start jobs from an offset (use if search failed and want to continue)
+    # proxy="http://jobspy:5a4vpWtj8EeJ2hoYzk@ca.smartproxy.com:20001",
+)
+
+# formatting for pandas
+pd.set_option('display.max_columns', None)
+pd.set_option('display.max_rows', None)
+pd.set_option('display.width', None)
+pd.set_option('display.max_colwidth', 50)  # set to 0 to see full job url / desc
+
+# 1: output to console
+print(jobs)
+
+# 2: output to .csv
+jobs.to_csv('./jobs.csv', index=False)
+print('outputted to jobs.csv')
+
+# 3: output to .xlsx
+# jobs.to_xlsx('jobs.xlsx', index=False)
+
+# 4: display in Jupyter Notebook (1. pip install jupyter 2. jupyter notebook)
+# display(jobs)
--- a/poetry.lock
+++ b/poetry.lock
--- a/pyproject.toml
+++ b/pyproject.toml
@@ -1,8 +1,9 @@
 [tool.poetry]
 name = "python-jobspy"
-version = "1.0.1"
+version = "1.1.8"
 description = "Job scraper for LinkedIn, Indeed & ZipRecruiter"
 authors = ["Zachary Hampton <zachary@zacharysproducts.com>", "Cullen Watson <cullen@cullen.ai>"]
+homepage = "https://github.com/cullenwatson/JobSpy"
 readme = "README.md"

 packages = [
@@ -24,4 +25,4 @@ jupyter = "^1.0.0"

 [build-system]
 requires = ["poetry-core"]
-build-backend = "poetry.core.masonry.api"
+build-backend = "poetry.core.masonry.api"
--- a/src/jobspy/core/init.py
+++ b/src/jobspy/core/init.py
--- a/src/jobspy/init.py
+++ b/src/jobspy/init.py
@@ -1,17 +1,19 @@
 import pandas as pd
-from typing import List, Tuple
+import concurrent.futures
+from concurrent.futures import ThreadPoolExecutor
+from typing import List, Tuple, Optional

-from .jobs import JobType
+from .jobs import JobType, Location
 from .scrapers.indeed import IndeedScraper
 from .scrapers.ziprecruiter import ZipRecruiterScraper
 from .scrapers.linkedin import LinkedInScraper
-from .scrapers import (
-    ScraperInput,
-    Site,
-    JobResponse,
+from .scrapers import ScraperInput, Site, JobResponse, Country
+from .scrapers.exceptions import (
+    LinkedInException,
+    IndeedException,
+    ZipRecruiterException,
 )

-
 SCRAPER_MAPPING = {
    Site.LINKEDIN: LinkedInScraper,
    Site.INDEED: IndeedScraper,
@@ -24,27 +26,45 @@ def _map_str_to_site(site_name: str) -> Site:


 def scrape_jobs(
-        site_name: str | Site | List[Site],
+        site_name: str | List[str] | Site | List[Site],
        search_term: str,
-
        location: str = "",
        distance: int = None,
        is_remote: bool = False,
-        job_type: JobType = None,
+        job_type: str = None,
        easy_apply: bool = False,  # linkedin
-        results_wanted: int = 15
+        results_wanted: int = 15,
+        country_indeed: str = "usa",
+        hyperlinks: bool = False,
+        proxy: Optional[str] = None,
+        offset: Optional[int] = 0
 ) -> pd.DataFrame:
    """
-    Asynchronously scrapes job data from multiple job sites.
+    Simultaneously scrapes job data from multiple job sites.
    :return: results_wanted: pandas dataframe containing job data
    """

-    if type(site_name) == str:
-        site_name = _map_str_to_site(site_name)
+    def get_enum_from_value(value_str):
+        for job_type in JobType:
+            if value_str in job_type.value:
+                return job_type
+        raise Exception(f"Invalid job type: {value_str}")
+
+    job_type = get_enum_from_value(job_type) if job_type else None
+
+    if type(site_name) == str:
+        site_type = [_map_str_to_site(site_name)]
+    else:  #: if type(site_name) == list
+        site_type = [
+            _map_str_to_site(site) if type(site) == str else site_name
+            for site in site_name
+        ]
+
+    country_enum = Country.from_string(country_indeed)

-    site_type = [site_name] if type(site_name) == Site else site_name
    scraper_input = ScraperInput(
        site_type=site_type,
+        country=country_enum,
        search_term=search_term,
        location=location,
        distance=distance,
@@ -52,67 +72,101 @@ def scrape_jobs(
        job_type=job_type,
        easy_apply=easy_apply,
        results_wanted=results_wanted,
+        offset=offset
    )

    def scrape_site(site: Site) -> Tuple[str, JobResponse]:
        scraper_class = SCRAPER_MAPPING[site]
-        scraper = scraper_class()
-        scraped_data: JobResponse = scraper.scrape(scraper_input)
+        scraper = scraper_class(proxy=proxy)

+        try:
+            scraped_data: JobResponse = scraper.scrape(scraper_input)
+        except (LinkedInException, IndeedException, ZipRecruiterException) as lie:
+            raise lie
+        except Exception as e:
+            # unhandled exceptions
+            if site == Site.LINKEDIN:
+                raise LinkedInException()
+            if site == Site.INDEED:
+                raise IndeedException()
+            if site == Site.ZIP_RECRUITER:
+                raise ZipRecruiterException()
+            else:
+                raise e
        return site.value, scraped_data

-    results = {}
-    for site in scraper_input.site_type:
+    site_to_jobs_dict = {}
+
+    def worker(site):
        site_value, scraped_data = scrape_site(site)
-        results[site_value] = scraped_data
+        return site_value, scraped_data

-    dfs = []
+    with ThreadPoolExecutor() as executor:
+        future_to_site = {
+            executor.submit(worker, site): site for site in scraper_input.site_type
+        }

-    for site, job_response in results.items():
+        for future in concurrent.futures.as_completed(future_to_site):
+            site_value, scraped_data = future.result()
+            site_to_jobs_dict[site_value] = scraped_data
+
+    jobs_dfs: List[pd.DataFrame] = []
+
+    for site, job_response in site_to_jobs_dict.items():
        for job in job_response.jobs:
-            data = job.dict()
-            data['site'] = site
-
-            # Formatting JobType
-            data['job_type'] = data['job_type'].value if data['job_type'] else None
-
-            # Formatting Location
-            location_obj = data.get('location')
-            if location_obj and isinstance(location_obj, dict):
-                data['city'] = location_obj.get('city', '')
-                data['state'] = location_obj.get('state', '')
-                data['country'] = location_obj.get('country', 'USA')
+            job_data = job.dict()
+            job_data[
+                "job_url_hyper"
+            ] = f'<a href="{job_data["job_url"]}">{job_data["job_url"]}</a>'
+            job_data["site"] = site
+            job_data["company"] = job_data["company_name"]
+            if job_data["job_type"]:
+                # Take the first value from the job type tuple
+                job_data["job_type"] = job_data["job_type"].value[0]
            else:
-                data['city'] = None
-                data['state'] = None
-                data['country'] = None
+                job_data["job_type"] = None

-            # Formatting Compensation
-            compensation_obj = data.get('compensation')
+            job_data["location"] = Location(**job_data["location"]).display_location()
+
+            compensation_obj = job_data.get("compensation")
            if compensation_obj and isinstance(compensation_obj, dict):
-                data['interval'] = compensation_obj.get('interval').value if compensation_obj.get('interval') else None
-                data['min_amount'] = compensation_obj.get('min_amount')
-                data['max_amount'] = compensation_obj.get('max_amount')
-                data['currency'] = compensation_obj.get('currency', 'USD')
+                job_data["interval"] = (
+                    compensation_obj.get("interval").value
+                    if compensation_obj.get("interval")
+                    else None
+                )
+                job_data["min_amount"] = compensation_obj.get("min_amount")
+                job_data["max_amount"] = compensation_obj.get("max_amount")
+                job_data["currency"] = compensation_obj.get("currency", "USD")
            else:
-                data['interval'] = None
-                data['min_amount'] = None
-                data['max_amount'] = None
-                data['currency'] = None
+                job_data["interval"] = None
+                job_data["min_amount"] = None
+                job_data["max_amount"] = None
+                job_data["currency"] = None

-            job_df = pd.DataFrame([data])
-            dfs.append(job_df)
+            job_df = pd.DataFrame([job_data])
+            jobs_dfs.append(job_df)

-    if dfs:
-        df = pd.concat(dfs, ignore_index=True)
-        desired_order = ['site', 'title', 'company_name', 'city', 'state','job_type',
-                         'interval', 'min_amount', 'max_amount',  'job_url', 'description',]
-        df = df[desired_order]
+    if jobs_dfs:
+        jobs_df = pd.concat(jobs_dfs, ignore_index=True)
+        desired_order: List[str] = [
+            "job_url_hyper" if hyperlinks else "job_url",
+            "site",
+            "title",
+            "company",
+            "location",
+            "job_type",
+            "date_posted",
+            "interval",
+            "benefits",
+            "min_amount",
+            "max_amount",
+            "currency",
+            "emails",
+            "description",
+        ]
+        jobs_formatted_df = jobs_df[desired_order]
    else:
-        df = pd.DataFrame()
-
-    return df
-
-
-
+        jobs_formatted_df = pd.DataFrame()

+    return jobs_formatted_df
--- a/src/jobspy/jobs/init.py
+++ b/src/jobspy/jobs/init.py
@@ -6,25 +6,160 @@ from pydantic import BaseModel, validator


 class JobType(Enum):
-    FULL_TIME = "fulltime"
-    PART_TIME = "parttime"
-    CONTRACT = "contract"
-    TEMPORARY = "temporary"
-    INTERNSHIP = "internship"
+    FULL_TIME = (
+        "fulltime",
+        "períodointegral",
+        "estágio/trainee",
+        "cunormăîntreagă",
+        "tiempocompleto",
+        "vollzeit",
+        "voltijds",
+        "tempointegral",
+        "全职",
+        "plnýúvazek",
+        "fuldtid",
+        "دوامكامل",
+        "kokopäivätyö",
+        "tempsplein",
+        "vollzeit",
+        "πλήρηςαπασχόληση",
+        "teljesmunkaidő",
+        "tempopieno",
+        "tempsplein",
+        "heltid",
+        "jornadacompleta",
+        "pełnyetat",
+        "정규직",
+        "100%",
+        "全職",
+        "งานประจำ",
+        "tamzamanlı",
+        "повназайнятість",
+        "toànthờigian",
+    )
+    PART_TIME = ("parttime", "teilzeit")
+    CONTRACT = ("contract", "contractor")
+    TEMPORARY = ("temporary",)
+    INTERNSHIP = ("internship", "prácticas", "ojt(onthejobtraining)", "praktikum")

-    PER_DIEM = "perdiem"
-    NIGHTS = "nights"
-    OTHER = "other"
-    SUMMER = "summer"
-    VOLUNTEER = "volunteer"
+    PER_DIEM = ("perdiem",)
+    NIGHTS = ("nights",)
+    OTHER = ("other",)
+    SUMMER = ("summer",)
+    VOLUNTEER = ("volunteer",)


+class Country(Enum):
+    ARGENTINA = ("argentina", "ar")
+    AUSTRALIA = ("australia", "au")
+    AUSTRIA = ("austria", "at")
+    BAHRAIN = ("bahrain", "bh")
+    BELGIUM = ("belgium", "be")
+    BRAZIL = ("brazil", "br")
+    CANADA = ("canada", "ca")
+    CHILE = ("chile", "cl")
+    CHINA = ("china", "cn")
+    COLOMBIA = ("colombia", "co")
+    COSTARICA = ("costa rica", "cr")
+    CZECHREPUBLIC = ("czech republic", "cz")
+    DENMARK = ("denmark", "dk")
+    ECUADOR = ("ecuador", "ec")
+    EGYPT = ("egypt", "eg")
+    FINLAND = ("finland", "fi")
+    FRANCE = ("france", "fr")
+    GERMANY = ("germany", "de")
+    GREECE = ("greece", "gr")
+    HONGKONG = ("hong kong", "hk")
+    HUNGARY = ("hungary", "hu")
+    INDIA = ("india", "in")
+    INDONESIA = ("indonesia", "id")
+    IRELAND = ("ireland", "ie")
+    ISRAEL = ("israel", "il")
+    ITALY = ("italy", "it")
+    JAPAN = ("japan", "jp")
+    KUWAIT = ("kuwait", "kw")
+    LUXEMBOURG = ("luxembourg", "lu")
+    MALAYSIA = ("malaysia", "malaysia")
+    MEXICO = ("mexico", "mx")
+    MOROCCO = ("morocco", "ma")
+    NETHERLANDS = ("netherlands", "nl")
+    NEWZEALAND = ("new zealand", "nz")
+    NIGERIA = ("nigeria", "ng")
+    NORWAY = ("norway", "no")
+    OMAN = ("oman", "om")
+    PAKISTAN = ("pakistan", "pk")
+    PANAMA = ("panama", "pa")
+    PERU = ("peru", "pe")
+    PHILIPPINES = ("philippines", "ph")
+    POLAND = ("poland", "pl")
+    PORTUGAL = ("portugal", "pt")
+    QATAR = ("qatar", "qa")
+    ROMANIA = ("romania", "ro")
+    SAUDIARABIA = ("saudi arabia", "sa")
+    SINGAPORE = ("singapore", "sg")
+    SOUTHAFRICA = ("south africa", "za")
+    SOUTHKOREA = ("south korea", "kr")
+    SPAIN = ("spain", "es")
+    SWEDEN = ("sweden", "se")
+    SWITZERLAND = ("switzerland", "ch")
+    TAIWAN = ("taiwan", "tw")
+    THAILAND = ("thailand", "th")
+    TURKEY = ("turkey", "tr")
+    UKRAINE = ("ukraine", "ua")
+    UNITEDARABEMIRATES = ("united arab emirates", "ae")
+    UK = ("uk", "uk")
+    USA = ("usa", "www")
+    URUGUAY = ("uruguay", "uy")
+    VENEZUELA = ("venezuela", "ve")
+    VIETNAM = ("vietnam", "vn")
+
+    # internal for ziprecruiter
+    US_CANADA = ("usa/ca", "www")
+
+    # internal for linkeind
+    WORLDWIDE = ("worldwide", "www")
+
+    def __new__(cls, country, domain):
+        obj = object.__new__(cls)
+        obj._value_ = country
+        obj.domain = domain
+        return obj
+
+    @property
+    def domain_value(self):
+        return self.domain
+
+    @classmethod
+    def from_string(cls, country_str: str):
+        """Convert a string to the corresponding Country enum."""
+        country_str = country_str.strip().lower()
+        for country in cls:
+            if country.value == country_str:
+                return country
+        valid_countries = [country.value for country in cls]
+        raise ValueError(
+            f"Invalid country string: '{country_str}'. Valid countries (only include this param for Indeed) are: {', '.join(valid_countries)}"
+        )
+

 class Location(BaseModel):
-    country: str = "USA"
-    city: str = None
+    country: Country = None
+    city: Optional[str] = None
    state: Optional[str] = None

+    def display_location(self) -> str:
+        location_parts = []
+        if self.city:
+            location_parts.append(self.city)
+        if self.state:
+            location_parts.append(self.state)
+        if self.country and self.country not in (Country.US_CANADA, Country.WORLDWIDE):
+            if self.country.value in ("usa", "uk"):
+                location_parts.append(self.country.value.upper())
+            else:
+                location_parts.append(self.country.value.title())
+        return ", ".join(location_parts)
+

 class CompensationInterval(Enum):
    YEARLY = "yearly"
@@ -35,10 +170,10 @@ class CompensationInterval(Enum):


 class Compensation(BaseModel):
-    interval: CompensationInterval
+    interval: Optional[CompensationInterval] = None
    min_amount: int = None
    max_amount: int = None
-    currency: str = "USD"
+    currency: Optional[str] = "USD"


 class JobPost(BaseModel):
@@ -47,29 +182,13 @@ class JobPost(BaseModel):
    job_url: str
    location: Optional[Location]

-    description: str = None
+    description: Optional[str] = None
    job_type: Optional[JobType] = None
    compensation: Optional[Compensation] = None
-    date_posted: date = None
+    date_posted: Optional[date] = None
+    benefits: Optional[str] = None
+    emails: Optional[list[str]] = None


 class JobResponse(BaseModel):
-    success: bool
-    error: str = None
-
-    total_results: Optional[int] = None
-
    jobs: list[JobPost] = []
-
-    returned_results: int = None
-
-    @validator("returned_results", pre=True, always=True)
-    def set_returned_results(cls, v, values):
-        jobs_list = values.get("jobs")
-
-        if v is None:
-            if jobs_list is not None:
-                return len(jobs_list)
-            else:
-                return 0
-        return v
--- a/src/jobspy/scrapers/init.py
+++ b/src/jobspy/scrapers/init.py
@@ -1,10 +1,5 @@
-from ..jobs import Enum, BaseModel, JobType, JobResponse
-from typing import List, Dict, Optional, Any
-
-
-class StatusException(Exception):
-    def __init__(self, status_code: int):
-        self.status_code = status_code
+from ..jobs import Enum, BaseModel, JobType, JobResponse, Country
+from typing import List, Optional, Any


 class Site(Enum):
@@ -18,26 +13,20 @@ class ScraperInput(BaseModel):
    search_term: str

    location: str = None
+    country: Optional[Country] = Country.USA
    distance: Optional[int] = None
    is_remote: bool = False
    job_type: Optional[JobType] = None
    easy_apply: bool = None  # linkedin
+    offset: int = 0

    results_wanted: int = 15


-class CommonResponse(BaseModel):
-    status: Optional[str]
-    error: Optional[str]
-    linkedin: Optional[Any] = None
-    indeed: Optional[Any] = None
-    zip_recruiter: Optional[Any] = None
-
-
 class Scraper:
-    def __init__(self, site: Site, url: str):
+    def __init__(self, site: Site, proxy: Optional[List[str]] = None):
        self.site = site
-        self.url = url
+        self.proxy = (lambda p: {"http": p, "https": p} if p else None)(proxy)

    def scrape(self, scraper_input: ScraperInput) -> JobResponse:
        ...
--- a/src/jobspy/scrapers/exceptions.py
+++ b/src/jobspy/scrapers/exceptions.py
@@ -0,0 +1,18 @@
+"""
+jobspy.scrapers.exceptions
+~~~~~~~~~~~~~~~~~~~
+
+This module contains the set of Scrapers' exceptions.
+"""
+
+
+class LinkedInException(Exception):
+    """Failed to scrape LinkedIn"""
+
+
+class IndeedException(Exception):
+    """Failed to scrape Indeed"""
+
+
+class ZipRecruiterException(Exception):
+    """Failed to scrape ZipRecruiter"""
--- a/src/jobspy/scrapers/indeed/init.py
+++ b/src/jobspy/scrapers/indeed/init.py
@@ -1,9 +1,15 @@
+"""
+jobspy.scrapers.indeed
+~~~~~~~~~~~~~~~~~~~
+
+This module contains routines to scrape Indeed.
+"""
 import re
-import sys
 import math
+import io
 import json
 from datetime import datetime
-from typing import Optional, Tuple, List
+from typing import Optional

 import tls_client
 import urllib.parse
@@ -11,28 +17,34 @@ from bs4 import BeautifulSoup
 from bs4.element import Tag
 from concurrent.futures import ThreadPoolExecutor, Future

-from ...jobs import JobPost, Compensation, CompensationInterval, Location, JobResponse, JobType
-from .. import Scraper, ScraperInput, Site, StatusException
-
-
-class ParsingException(Exception):
-    pass
+from ..exceptions import IndeedException
+from ...jobs import (
+    JobPost,
+    Compensation,
+    CompensationInterval,
+    Location,
+    JobResponse,
+    JobType,
+)
+from .. import Scraper, ScraperInput, Site
+from ...utils import extract_emails_from_text


 class IndeedScraper(Scraper):
-    def __init__(self):
+    def __init__(self, proxy: Optional[str] = None):
        """
        Initializes IndeedScraper with the Indeed job search url
        """
+        self.url = None
+        self.country = None
        site = Site(Site.INDEED)
-        url = "https://www.indeed.com"
-        super().__init__(site, url)
+        super().__init__(site, proxy=proxy)

        self.jobs_per_page = 15
        self.seen_urls = set()

    def scrape_page(
-        self, scraper_input: ScraperInput, page: int, session: tls_client.Session
+            self, scraper_input: ScraperInput, page: int, session: tls_client.Session
    ) -> tuple[list[JobPost], int]:
        """
        Scrapes a page of Indeed for jobs with scraper_input criteria
@@ -41,16 +53,21 @@ class IndeedScraper(Scraper):
        :param session:
        :return: jobs found on page, total number of jobs found for search
        """
+        self.country = scraper_input.country
+        domain = self.country.domain_value
+        self.url = f"https://{domain}.indeed.com"

-        job_list = []
+        job_list: list[JobPost] = []

        params = {
            "q": scraper_input.search_term,
            "l": scraper_input.location,
-            "radius": scraper_input.distance,
            "filter": 0,
-            "start": 0 + page * 10,
+            "start": scraper_input.offset + page * 10,
        }
+        if scraper_input.distance:
+            params["radius"] = scraper_input.distance
+
        sc_values = []
        if scraper_input.is_remote:
            sc_values.append("attr(DSQF7)")
@@ -59,17 +76,26 @@ class IndeedScraper(Scraper):

        if sc_values:
            params["sc"] = "0kf:" + "".join(sc_values) + ";"
-        response = session.get(self.url + "/jobs", params=params)
-
-        if (
-            response.status_code != 200
-            and response.status_code != 307
-        ):
-            raise StatusException(response.status_code)
+        try:
+            response = session.get(
+                f"{self.url}/jobs",
+                params=params,
+                allow_redirects=True,
+                proxy=self.proxy,
+                timeout_seconds=10,
+            )
+            if response.status_code not in range(200, 400):
+                raise IndeedException(
+                    f"bad response with status code: {response.status_code}"
+                )
+        except Exception as e:
+            if "Proxy responded with" in str(e):
+                raise IndeedException("bad proxy")
+            raise IndeedException(str(e))

        soup = BeautifulSoup(response.content, "html.parser")
-        if "did not match any jobs" in str(soup):
-            raise ParsingException("Search did not match any jobs")
+        if "did not match any jobs" in response.text:
+            raise IndeedException("Parsing exception: Search did not match any jobs")

        jobs = IndeedScraper.parse_jobs(
            soup
@@ -77,11 +103,11 @@ class IndeedScraper(Scraper):
        total_num_jobs = IndeedScraper.total_jobs(soup)

        if (
-            not jobs.get("metaData", {})
-            .get("mosaicProviderJobCardsModel", {})
-            .get("results")
+                not jobs.get("metaData", {})
+                        .get("mosaicProviderJobCardsModel", {})
+                        .get("results")
        ):
-            raise Exception("No jobs found.")
+            raise IndeedException("No jobs found.")

        def process_job(job) -> Optional[JobPost]:
            job_url = f'{self.url}/jobs/viewjob?jk={job["jobkey"]}'
@@ -89,8 +115,6 @@ class IndeedScraper(Scraper):
            if job_url in self.seen_urls:
                return None

-            snippet_html = BeautifulSoup(job["snippet"], "html.parser")
-
            extracted_salary = job.get("extractedSalary")
            compensation = None
            if extracted_salary:
@@ -115,11 +139,13 @@ class IndeedScraper(Scraper):
            date_posted = date_posted.strftime("%Y-%m-%d")

            description = self.get_description(job_url, session)
-            li_elements = snippet_html.find_all("li")
-            if description is None and li_elements:
-                description = " ".join(li.text for li in li_elements)
+            emails = extract_emails_from_text(description)
+            with io.StringIO(job["snippet"]) as f:
+                soup_io = BeautifulSoup(f, "html.parser")
+                li_elements = soup_io.find_all("li")
+                if description is None and li_elements:
+                    description = " ".join(li.text for li in li_elements)

-            first_li = snippet_html.find("li")
            job_post = JobPost(
                title=job["normTitle"],
                description=description,
@@ -127,7 +153,9 @@ class IndeedScraper(Scraper):
                location=Location(
                    city=job.get("jobLocationCity"),
                    state=job.get("jobLocationState"),
+                    country=self.country,
                ),
+                emails=extract_emails_from_text(description),
                job_type=job_type,
                compensation=compensation,
                date_posted=date_posted,
@@ -135,9 +163,11 @@ class IndeedScraper(Scraper):
            )
            return job_post

-        with ThreadPoolExecutor(max_workers=10) as executor:
-            job_results: list[Future] = [executor.submit(process_job, job) for job in
-                                         jobs["metaData"]["mosaicProviderJobCardsModel"]["results"]]
+        with ThreadPoolExecutor(max_workers=1) as executor:
+            job_results: list[Future] = [
+                executor.submit(process_job, job)
+                for job in jobs["metaData"]["mosaicProviderJobCardsModel"]["results"]
+            ]

        job_list = [result.result() for result in job_results if result.result()]

@@ -154,51 +184,33 @@ class IndeedScraper(Scraper):
        )

        pages_to_process = (
-            math.ceil(scraper_input.results_wanted / self.jobs_per_page) - 1
+                math.ceil(scraper_input.results_wanted / self.jobs_per_page) - 1
        )

-        try:
-            #: get first page to initialize session
-            job_list, total_results = self.scrape_page(scraper_input, 0, session)
+        #: get first page to initialize session
+        job_list, total_results = self.scrape_page(scraper_input, 0, session)

-            with ThreadPoolExecutor(max_workers=10) as executor:
-                futures: list[Future] = [
-                    executor.submit(self.scrape_page, scraper_input, page, session)
-                    for page in range(1, pages_to_process + 1)
-                ]
+        with ThreadPoolExecutor(max_workers=1) as executor:
+            futures: list[Future] = [
+                executor.submit(self.scrape_page, scraper_input, page, session)
+                for page in range(1, pages_to_process + 1)
+            ]

-                for future in futures:
-                    jobs, _ = future.result()
+            for future in futures:
+                jobs, _ = future.result()

-                    job_list += jobs
-        except StatusException as e:
-            return JobResponse(
-                success=False,
-                error=f"Indeed returned status code {e.status_code}",
-            )
-
-        except ParsingException as e:
-            return JobResponse(
-                success=False,
-                error=f"Indeed failed to parse response: {e}",
-            )
-        except Exception as e:
-            return JobResponse(
-                success=False,
-                error=f"Indeed failed to scrape: {e}",
-            )
+                job_list += jobs

        if len(job_list) > scraper_input.results_wanted:
            job_list = job_list[: scraper_input.results_wanted]

        job_response = JobResponse(
-            success=True,
            jobs=job_list,
            total_results=total_results,
        )
        return job_response

-    def get_description(self, job_page_url: str, session: tls_client.Session) -> str:
+    def get_description(self, job_page_url: str, session: tls_client.Session) -> Optional[str]:
        """
        Retrieves job description by going to the job page url
        :param job_page_url:
@@ -210,7 +222,12 @@ class IndeedScraper(Scraper):
        jk_value = params.get("jk", [None])[0]
        formatted_url = f"{self.url}/viewjob?jk={jk_value}&spa=1"

-        response = session.get(formatted_url, allow_redirects=True)
+        try:
+            response = session.get(
+                formatted_url, allow_redirects=True, timeout_seconds=5, proxy=self.proxy
+            )
+        except Exception as e:
+            return None

        if response.status_code not in range(200, 400):
            return None
@@ -218,9 +235,10 @@ class IndeedScraper(Scraper):
        raw_description = response.json()["body"]["jobInfoWrapperModel"][
            "jobInfoModel"
        ]["sanitizedJobDescription"]
-        soup = BeautifulSoup(raw_description, "html.parser")
-        text_content = " ".join(soup.get_text().split()).strip()
-        return text_content
+        with io.StringIO(raw_description) as f:
+            soup = BeautifulSoup(f, "html.parser")
+            text_content = " ".join(soup.get_text().split()).strip()
+            return text_content

    @staticmethod
    def get_job_type(job: dict) -> Optional[JobType]:
@@ -232,13 +250,21 @@ class IndeedScraper(Scraper):
        for taxonomy in job["taxonomyAttributes"]:
            if taxonomy["label"] == "job-types":
                if len(taxonomy["attributes"]) > 0:
-                    job_type_str = (
-                        taxonomy["attributes"][0]["label"]
-                        .replace("-", "_")
-                        .replace(" ", "_")
-                        .upper()
-                    )
-                    return JobType[job_type_str]
+                    label = taxonomy["attributes"][0].get("label")
+                    if label:
+                        job_type_str = label.replace("-", "").replace(" ", "").lower()
+                        return IndeedScraper.get_enum_from_job_type(job_type_str)
+        return None
+
+    @staticmethod
+    def get_enum_from_job_type(job_type_str):
+        """
+        Given a string, returns the corresponding JobType enum member if a match is found.
+        for job_type in JobType:
+        """
+        for job_type in JobType:
+            if job_type_str in job_type.value:
+                return job_type
        return None

    @staticmethod
@@ -258,9 +284,9 @@ class IndeedScraper(Scraper):

            for tag in script_tags:
                if (
-                    tag.string
-                    and "mosaic.providerData" in tag.string
-                    and "mosaic-provider-jobcards" in tag.string
+                        tag.string
+                        and "mosaic.providerData" in tag.string
+                        and "mosaic-provider-jobcards" in tag.string
                ):
                    return tag
            return None
@@ -276,9 +302,9 @@ class IndeedScraper(Scraper):
                jobs = json.loads(m.group(1).strip())
                return jobs
            else:
-                raise ParsingException("Could not find mosaic provider job cards data")
+                raise IndeedException("Could not find mosaic provider job cards data")
        else:
-            raise ParsingException(
+            raise IndeedException(
                "Could not find a script tag containing mosaic provider data"
            )

@@ -289,7 +315,7 @@ class IndeedScraper(Scraper):
        :param soup:
        :return: total_num_jobs
        """
-        script = soup.find("script", string=lambda t: "window._initialData" in t)
+        script = soup.find("script", string=lambda t: t and "window._initialData" in t)

        pattern = re.compile(r"window._initialData\s*=\s*({.*})\s*;", re.DOTALL)
        match = pattern.search(script.string)
--- a/src/jobspy/scrapers/linkedin/init.py
+++ b/src/jobspy/scrapers/linkedin/init.py
@@ -1,22 +1,43 @@
-from typing import Optional, Tuple
+"""
+jobspy.scrapers.linkedin
+~~~~~~~~~~~~~~~~~~~
+
+This module contains routines to scrape LinkedIn.
+"""
+from typing import Optional
 from datetime import datetime

 import requests
+import time
+from requests.exceptions import ProxyError
+from concurrent.futures import ThreadPoolExecutor, as_completed
 from bs4 import BeautifulSoup
 from bs4.element import Tag
+from threading import Lock

 from .. import Scraper, ScraperInput, Site
-from ...jobs import JobPost, Location, JobResponse, JobType, Compensation, CompensationInterval
+from ..exceptions import LinkedInException
+from ...jobs import (
+    JobPost,
+    Location,
+    JobResponse,
+    JobType,
+)
+from ...utils import extract_emails_from_text


 class LinkedInScraper(Scraper):
-    def __init__(self):
+    MAX_RETRIES = 3
+    DELAY = 10
+
+    def __init__(self, proxy: Optional[str] = None):
        """
        Initializes LinkedInScraper with the LinkedIn job search url
        """
        site = Site(Site.LINKEDIN)
-        url = "https://www.linkedin.com"
-        super().__init__(site, url)
+        self.country = "worldwide"
+        self.url = "https://www.linkedin.com"
+        super().__init__(site, proxy=proxy)

    def scrape(self, scraper_input: ScraperInput) -> JobResponse:
        """
@@ -26,9 +47,10 @@ class LinkedInScraper(Scraper):
        """
        job_list: list[JobPost] = []
        seen_urls = set()
-        page, processed_jobs, job_count = 0, 0, 0
+        url_lock = Lock()
+        page = scraper_input.offset // 25 + 25 if scraper_input.offset else 0

-        def job_type_code(job_type):
+        def job_type_code(job_type_enum):
            mapping = {
                JobType.FULL_TIME: "F",
                JobType.PART_TIME: "P",
@@ -37,119 +59,134 @@ class LinkedInScraper(Scraper):
                JobType.TEMPORARY: "T",
            }

-            return mapping.get(job_type, "")
+            return mapping.get(job_type_enum, "")

-        with requests.Session() as session:
-            while len(job_list) < scraper_input.results_wanted:
-                params = {
-                    "keywords": scraper_input.search_term,
-                    "location": scraper_input.location,
-                    "distance": scraper_input.distance,
-                    "f_WT": 2 if scraper_input.is_remote else None,
-                    "f_JT": job_type_code(scraper_input.job_type)
-                    if scraper_input.job_type
-                    else None,
-                    "pageNum": page,
-                    "f_AL": "true" if scraper_input.easy_apply else None,
-                }
+        while len(job_list) < scraper_input.results_wanted and page < 1000:
+            params = {
+                "keywords": scraper_input.search_term,
+                "location": scraper_input.location,
+                "distance": scraper_input.distance,
+                "f_WT": 2 if scraper_input.is_remote else None,
+                "f_JT": job_type_code(scraper_input.job_type)
+                if scraper_input.job_type
+                else None,
+                "pageNum": 0,
+                page: page + scraper_input.offset,
+                "f_AL": "true" if scraper_input.easy_apply else None,
+            }

-                params = {k: v for k, v in params.items() if v is not None}
-                response = session.get(
-                    f"{self.url}/jobs/search", params=params, allow_redirects=True
-                )
+            params = {k: v for k, v in params.items() if v is not None}

-                if response.status_code != 200:
-                    return JobResponse(
-                        success=False,
-                        error=f"Response returned {response.status_code}",
+            params = {k: v for k, v in params.items() if v is not None}
+            retries = 0
+            while retries < self.MAX_RETRIES:
+                try:
+                    response = requests.get(
+                        f"{self.url}/jobs-guest/jobs/api/seeMoreJobPostings/search?",
+                        params=params,
+                        allow_redirects=True,
+                        proxies=self.proxy,
+                        timeout=10,
                    )
+                    response.raise_for_status()

-                soup = BeautifulSoup(response.text, "html.parser")
-
-                if page == 0:
-                    job_count_text = soup.find(
-                        "span", class_="results-context-header__job-count"
-                    ).text
-                    job_count = int("".join(filter(str.isdigit, job_count_text)))
-
-                for job_card in soup.find_all(
-                    "div",
-                    class_="base-card relative w-full hover:no-underline focus:no-underline base-card--link base-search-card base-search-card--link job-search-card",
-                ):
-                    processed_jobs += 1
-                    data_entity_urn = job_card.get("data-entity-urn", "")
-                    job_id = (
-                        data_entity_urn.split(":")[-1] if data_entity_urn else "N/A"
-                    )
-                    job_url = f"{self.url}/jobs/view/{job_id}"
-                    if job_url in seen_urls:
-                        continue
-                    seen_urls.add(job_url)
-                    job_info = job_card.find("div", class_="base-search-card__info")
-                    if job_info is None:
-                        continue
-                    title_tag = job_info.find("h3", class_="base-search-card__title")
-                    title = title_tag.text.strip() if title_tag else "N/A"
-
-                    company_tag = job_info.find("a", class_="hidden-nested-link")
-                    company = company_tag.text.strip() if company_tag else "N/A"
-
-                    metadata_card = job_info.find(
-                        "div", class_="base-search-card__metadata"
-                    )
-                    location: Location = LinkedInScraper.get_location(metadata_card)
-
-                    datetime_tag = metadata_card.find(
-                        "time", class_="job-search-card__listdate"
-                    )
-                    description, job_type = LinkedInScraper.get_description(job_url)
-                    if datetime_tag:
-                        datetime_str = datetime_tag["datetime"]
-                        date_posted = datetime.strptime(datetime_str, "%Y-%m-%d")
-                    else:
-                        date_posted = None
-
-                    job_post = JobPost(
-                        title=title,
-                        description=description,
-                        company_name=company,
-                        location=location,
-                        date_posted=date_posted,
-                        job_url=job_url,
-                        job_type=job_type,
-                        compensation=Compensation(interval=CompensationInterval.YEARLY, currency="USD")
-                    )
-                    job_list.append(job_post)
-                    if (
-                        len(job_list) >= scraper_input.results_wanted
-                        or processed_jobs >= job_count
-                    ):
-                        break
-                if (
-                    len(job_list) >= scraper_input.results_wanted
-                    or processed_jobs >= job_count
-                ):
                    break
+                except requests.HTTPError as e:
+                    if hasattr(e, 'response') and e.response is not None:
+                        if e.response.status_code == 429:
+                            time.sleep(self.DELAY)
+                            retries += 1
+                            continue
+                        else:
+                            raise LinkedInException(f"bad response status code: {e.response.status_code}")
+                    else:
+                        raise
+                except ProxyError as e:
+                    raise LinkedInException("bad proxy")
+                except Exception as e:
+                    raise LinkedInException(str(e))
+            else:
+                # Raise an exception if the maximum number of retries is reached
+                raise LinkedInException("Max retries reached, failed to get a valid response")

-                page += 1
+            soup = BeautifulSoup(response.text, "html.parser")
+
+            with ThreadPoolExecutor(max_workers=5) as executor:
+                futures = []
+                for job_card in soup.find_all("div", class_="base-search-card"):
+                    job_url = None
+                    href_tag = job_card.find("a", class_="base-card__full-link")
+                    if href_tag and "href" in href_tag.attrs:
+                        href = href_tag.attrs["href"].split("?")[0]
+                        job_id = href.split("-")[-1]
+                        job_url = f"{self.url}/jobs/view/{job_id}"
+
+                    with url_lock:
+                        if job_url in seen_urls:
+                            continue
+                        seen_urls.add(job_url)
+
+                    futures.append(executor.submit(self.process_job, job_card, job_url))
+
+                for future in as_completed(futures):
+                    try:
+                        job_post = future.result()
+                        if job_post:
+                            job_list.append(job_post)
+                    except Exception as e:
+                        raise LinkedInException("Exception occurred while processing jobs")
+            page += 25

        job_list = job_list[: scraper_input.results_wanted]
-        job_response = JobResponse(
-            success=True,
-            jobs=job_list,
-            total_results=job_count,
-        )
-        return job_response
+        return JobResponse(jobs=job_list)

-    @staticmethod
-    def get_description(job_page_url: str) -> Optional[str]:
+    def process_job(self, job_card: Tag, job_url: str) -> Optional[JobPost]:
+        title_tag = job_card.find("span", class_="sr-only")
+        title = title_tag.get_text(strip=True) if title_tag else "N/A"
+
+        company_tag = job_card.find("h4", class_="base-search-card__subtitle")
+        company_a_tag = company_tag.find("a") if company_tag else None
+        company = company_a_tag.get_text(strip=True) if company_a_tag else "N/A"
+
+        metadata_card = job_card.find("div", class_="base-search-card__metadata")
+        location = self.get_location(metadata_card)
+
+        datetime_tag = metadata_card.find("time", class_="job-search-card__listdate") if metadata_card else None
+        date_posted = None
+        if datetime_tag and "datetime" in datetime_tag.attrs:
+            datetime_str = datetime_tag["datetime"]
+            try:
+                date_posted = datetime.strptime(datetime_str, "%Y-%m-%d")
+            except Exception as e:
+                date_posted = None
+        benefits_tag = job_card.find("span", class_="result-benefits__text")
+        benefits = " ".join(benefits_tag.get_text().split()) if benefits_tag else None
+
+        description, job_type = self.get_job_description(job_url)
+
+        return JobPost(
+            title=title,
+            description=description,
+            company_name=company,
+            location=location,
+            date_posted=date_posted,
+            job_url=job_url,
+            job_type=job_type,
+            benefits=benefits,
+            emails=extract_emails_from_text(description)
+        )
+
+    def get_job_description(self, job_page_url: str) -> tuple[None, None] | tuple[
+        str | None, tuple[str | None, JobType | None]]:
        """
        Retrieves job description by going to the job page url
        :param job_page_url:
        :return: description or None
        """
-        response = requests.get(job_page_url, allow_redirects=True)
-        if response.status_code not in range(200, 400):
+        try:
+            response = requests.get(job_page_url, timeout=5, proxies=self.proxy)
+            response.raise_for_status()
+        except Exception as e:
            return None, None

        soup = BeautifulSoup(response.text, "html.parser")
@@ -157,19 +194,19 @@ class LinkedInScraper(Scraper):
            "div", class_=lambda x: x and "show-more-less-html__markup" in x
        )

-        text_content = None
+        description = None
        if div_content:
-            text_content = " ".join(div_content.get_text().split()).strip()
+            description = " ".join(div_content.get_text().split()).strip()

        def get_job_type(
-            soup: BeautifulSoup,
-        ) -> Tuple[Optional[str], Optional[JobType]]:
+                soup_job_type: BeautifulSoup,
+        ) -> JobType | None:
            """
            Gets the job type from job page
-            :param soup:
+            :param soup_job_type:
            :return: JobType
            """
-            h3_tag = soup.find(
+            h3_tag = soup_job_type.find(
                "h3",
                class_="description__job-criteria-subheader",
                string=lambda text: "Employment type" in text,
@@ -186,17 +223,24 @@ class LinkedInScraper(Scraper):
                    employment_type = employment_type.lower()
                    employment_type = employment_type.replace("-", "")

-            return JobType(employment_type)
+            return LinkedInScraper.get_enum_from_value(employment_type)

-        return text_content, get_job_type(soup)
+        return description, get_job_type(soup)

    @staticmethod
-    def get_location(metadata_card: Optional[Tag]) -> Location:
+    def get_enum_from_value(value_str):
+        for job_type in JobType:
+            if value_str in job_type.value:
+                return job_type
+        return None
+
+    def get_location(self, metadata_card: Optional[Tag]) -> Location:
        """
        Extracts the location data from the job metadata card.
        :param metadata_card
        :return: location
        """
+        location = Location(country=self.country)
        if metadata_card is not None:
            location_tag = metadata_card.find(
                "span", class_="job-search-card__location"
@@ -208,6 +252,7 @@ class LinkedInScraper(Scraper):
                location = Location(
                    city=city,
                    state=state,
+                    country=self.country,
                )

        return location
--- a/src/jobspy/scrapers/ziprecruiter/init.py
+++ b/src/jobspy/scrapers/ziprecruiter/init.py
@@ -1,27 +1,44 @@
+"""
+jobspy.scrapers.ziprecruiter
+~~~~~~~~~~~~~~~~~~~
+
+This module contains routines to scrape ZipRecruiter.
+"""
 import math
 import json
 import re
-from datetime import datetime
-from typing import Optional, Tuple, List
-from urllib.parse import urlparse, parse_qs
+from datetime import datetime, date
+from typing import Optional, Tuple, Any
+from urllib.parse import urlparse, parse_qs, urlunparse

 import tls_client
+import requests
 from bs4 import BeautifulSoup
 from bs4.element import Tag
 from concurrent.futures import ThreadPoolExecutor, Future

-from .. import Scraper, ScraperInput, Site, StatusException
-from ...jobs import JobPost, Compensation, CompensationInterval, Location, JobResponse, JobType
+from .. import Scraper, ScraperInput, Site
+from ..exceptions import ZipRecruiterException
+from ...jobs import (
+    JobPost,
+    Compensation,
+    CompensationInterval,
+    Location,
+    JobResponse,
+    JobType,
+    Country,
+)
+from ...utils import extract_emails_from_text


 class ZipRecruiterScraper(Scraper):
-    def __init__(self):
+    def __init__(self, proxy: Optional[str] = None):
        """
        Initializes LinkedInScraper with the ZipRecruiter job search url
        """
        site = Site(Site.ZIP_RECRUITER)
-        url = "https://www.ziprecruiter.com"
-        super().__init__(site, url)
+        self.url = "https://www.ziprecruiter.com"
+        super().__init__(site, proxy=proxy)

        self.jobs_per_page = 20
        self.seen_urls = set()
@@ -29,76 +46,69 @@ class ZipRecruiterScraper(Scraper):
            client_identifier="chrome112", random_tls_extension_order=True
        )

-    def scrape_page(
-        self, scraper_input: ScraperInput, page: int
-    ) -> tuple[list[JobPost], int | None]:
+    def find_jobs_in_page(
+            self, scraper_input: ScraperInput, page: int
+    ) -> list[JobPost]:
        """
        Scrapes a page of ZipRecruiter for jobs with scraper_input criteria
        :param scraper_input:
        :param page:
-        :param session:
-        :return: jobs found on page, total number of jobs found for search
+        :return: jobs found on page
        """
-
-        job_list = []
-
-        job_type_value = None
-        if scraper_input.job_type:
-            if scraper_input.job_type.value == "fulltime":
-                job_type_value = "full_time"
-            elif scraper_input.job_type.value == "parttime":
-                job_type_value = "part_time"
-            else:
-                job_type_value = scraper_input.job_type.value
-
-        params = {
-            "search": scraper_input.search_term,
-            "location": scraper_input.location,
-            "page": page,
-            "form": "jobs-landing"
-        }
-
-        if scraper_input.is_remote:
-            params["refine_by_location_type"] = "only_remote"
-
-        if scraper_input.distance:
-            params["radius"] = scraper_input.distance
-
-        if job_type_value:
-            params["refine_by_employment"] = f"employment_type:employment_type:{job_type_value}"
-
-        response = self.session.get(
-            self.url + "/jobs-search",
-            headers=ZipRecruiterScraper.headers(),
-            params=params,
-        )
-
-        if response.status_code != 200:
-            raise StatusException(response.status_code)
-
-        html_string = response.text
-        soup = BeautifulSoup(html_string, "html.parser")
-
-        script_tag = soup.find("script", {"id": "js_variables"})
-        data = json.loads(script_tag.string)
-
-        if page == 1:
-            job_count = int(data["totalJobCount"].replace(",", ""))
+        job_list: list[JobPost] = []
+        try:
+            response = self.session.get(
+                f"{self.url}/jobs-search",
+                headers=ZipRecruiterScraper.headers(),
+                params=ZipRecruiterScraper.add_params(scraper_input, page),
+                allow_redirects=True,
+                proxy=self.proxy,
+                timeout_seconds=10,
+            )
+            if response.status_code != 200:
+                raise ZipRecruiterException(
+                    f"bad response status code: {response.status_code}"
+                )
+        except Exception as e:
+            if "Proxy responded with non 200 code" in str(e):
+                raise ZipRecruiterException("bad proxy")
+            raise ZipRecruiterException(str(e))
        else:
-            job_count = None
+            soup = BeautifulSoup(response.text, "html.parser")
+            js_tag = soup.find("script", {"id": "js_variables"})
+
+        if js_tag:
+            page_json = json.loads(js_tag.string)
+            jobs_list = page_json.get("jobList")
+            if jobs_list:
+                page_variant = "javascript"
+                # print('type javascript', len(jobs_list))
+            else:
+                page_variant = "html_2"
+                jobs_list = soup.find_all("div", {"class": "job_content"})
+                # print('type 2 html', len(jobs_list))
+        else:
+            page_variant = "html_1"
+            jobs_list = soup.find_all("li", {"class": "job-listing"})
+            # print('type 1 html', len(jobs_list))

        with ThreadPoolExecutor(max_workers=10) as executor:
-            if "jobList" in data and data["jobList"]:
-                jobs_js = data["jobList"]
-                job_results = [executor.submit(self.process_job_js, job) for job in jobs_js]
-            else:
-                jobs_html = soup.find_all("div", {"class": "job_content"})
-                job_results = [executor.submit(self.process_job_html, job) for job in
-                               jobs_html]
+            if page_variant == "javascript":
+                job_results = [
+                    executor.submit(self.process_job_javascript, job)
+                    for job in jobs_list
+                ]
+            elif page_variant == "html_1":
+                job_results = [
+                    executor.submit(self.process_job_html_1, job) for job in jobs_list
+                ]
+            elif page_variant == "html_2":
+                job_results = [
+                    executor.submit(self.process_job_html_2, job) for job in jobs_list
+                ]

        job_list = [result.result() for result in job_results if result.result()]
-
-        return job_list, job_count
+        return job_list

    def scrape(self, scraper_input: ScraperInput) -> JobResponse:
        """
@@ -106,84 +116,95 @@ class ZipRecruiterScraper(Scraper):
        :param scraper_input:
        :return: job_response
        """
-
-
-        pages_to_process = max(3, math.ceil(scraper_input.results_wanted / self.jobs_per_page))
-
-        try:
-            #: get first page to initialize session
-            job_list, total_results = self.scrape_page(scraper_input, 1)
-
-            with ThreadPoolExecutor(max_workers=10) as executor:
-                futures: list[Future] = [
-                    executor.submit(self.scrape_page, scraper_input, page)
-                    for page in range(2, pages_to_process + 1)
-                ]
-
-                for future in futures:
-                    jobs, _ = future.result()
-
-                    job_list += jobs
-
-
-        except StatusException as e:
-            return JobResponse(
-                success=False,
-                error=f"ZipRecruiter returned status code {e.status_code}",
-            )
-        except Exception as e:
-            return JobResponse(
-                success=False,
-                error=f"ZipRecruiter failed to scrape: {e}",
-            )
-
-        #: note: this does not handle if the results are more or less than the results_wanted
-
-        if len(job_list) > scraper_input.results_wanted:
-            job_list = job_list[: scraper_input.results_wanted]
-
-        job_response = JobResponse(
-            success=True,
-            jobs=job_list,
-            total_results=total_results,
+        start_page = (scraper_input.offset // self.jobs_per_page) + 1 if scraper_input.offset else 1
+        #: get first page to initialize session
+        job_list: list[JobPost] = self.find_jobs_in_page(scraper_input, start_page)
+        pages_to_process = max(
+            3, math.ceil(scraper_input.results_wanted / self.jobs_per_page)
        )
-        return job_response

-    def process_job_html(self, job: Tag) -> Optional[JobPost]:
+        with ThreadPoolExecutor(max_workers=10) as executor:
+            futures: list[Future] = [
+                executor.submit(self.find_jobs_in_page, scraper_input, page)
+                for page in range(start_page + 1, start_page + pages_to_process + 2)
+            ]
+
+            for future in futures:
+                jobs = future.result()
+
+                job_list += jobs
+
+        job_list = job_list[: scraper_input.results_wanted]
+        return JobResponse(jobs=job_list)
+
+    def process_job_html_1(self, job: Tag) -> Optional[JobPost]:
        """
        Parses a job from the job content tag
        :param job: BeautifulSoup Tag for one job post
        :return JobPost
+        TODO this method isnt finished due to not encountering this type of html often
        """
-        job_url = job.find("a", {"class": "job_link"})["href"]
+        job_url = self.cleanurl(job.find("a", {"class": "job_link"})["href"])
        if job_url in self.seen_urls:
            return None

        title = job.find("h2", {"class": "title"}).text
        company = job.find("a", {"class": "company_name"}).text.strip()

-        description, updated_job_url = self.get_description(
-            job_url
-        )
-        if updated_job_url is not None:
-            job_url = updated_job_url
+        description, updated_job_url = self.get_description(job_url)
+        # job_url = updated_job_url if updated_job_url else job_url
        if description is None:
            description = job.find("p", {"class": "job_snippet"}).text.strip()

        job_type_element = job.find("li", {"class": "perk_item perk_type"})
+        job_type = None
        if job_type_element:
            job_type_text = (
-                job_type_element.text.strip()
+                job_type_element.text.strip().lower().replace("_", "").replace(" ", "")
+            )
+            job_type = ZipRecruiterScraper.get_job_type_enum(job_type_text)
+
+        date_posted = ZipRecruiterScraper.get_date_posted(job)
+
+        job_post = JobPost(
+            title=title,
+            description=description,
+            company_name=company,
+            location=ZipRecruiterScraper.get_location(job),
+            job_type=job_type,
+            compensation=ZipRecruiterScraper.get_compensation(job),
+            date_posted=date_posted,
+            job_url=job_url,
+            emails=extract_emails_from_text(description),
+        )
+        return job_post
+
+    def process_job_html_2(self, job: Tag) -> Optional[JobPost]:
+        """
+        Parses a job from the job content tag for a second variat of HTML that ZR uses
+        :param job: BeautifulSoup Tag for one job post
+        :return JobPost
+        """
+        job_url = self.cleanurl(job.find("a", class_="job_link")["href"])
+        title = job.find("h2", class_="title").text
+        company = job.find("a", class_="company_name").text.strip()
+
+        description, updated_job_url = self.get_description(job_url)
+        # job_url = updated_job_url if updated_job_url else job_url
+        if description is None:
+            description = job.find("p", class_="job_snippet").get_text().strip()
+
+        job_type_text = job.find("li", class_="perk_item perk_type")
+        job_type = None
+        if job_type_text:
+            job_type_text = (
+                job_type_text.get_text()
+                .strip()
                .lower()
                .replace("-", "")
                .replace(" ", "")
            )
-            if job_type_text == "contractor":
-                job_type_text = "contract"
-            job_type = JobType(job_type_text)
-        else:
-            job_type = None
-
+            job_type = ZipRecruiterScraper.get_job_type_enum(job_type_text)
        date_posted = ZipRecruiterScraper.get_date_posted(job)

        job_post = JobPost(
@@ -198,31 +219,37 @@ class ZipRecruiterScraper(Scraper):
        )
        return job_post

-    def process_job_js(self, job: dict) -> JobPost:
-        # Map the job data to the expected fields by the Pydantic model
+    def process_job_javascript(self, job: dict) -> JobPost:
        title = job.get("Title")
-        description = BeautifulSoup(job.get("Snippet","").strip(), "html.parser").get_text()
+        job_url = self.cleanurl(job.get("JobURL"))
+
+        description, updated_job_url = self.get_description(job_url)
+        # job_url = updated_job_url if updated_job_url else job_url
+        if description is None:
+            description = BeautifulSoup(
+                job.get("Snippet", "").strip(), "html.parser"
+            ).get_text()

        company = job.get("OrgName")
-        location = Location(city=job.get("City"), state=job.get("State"))
-        try:
-            job_type = ZipRecruiterScraper.job_type_from_string(job.get("EmploymentType", "").replace("-", "_").lower())
-        except ValueError:
-            # print(f"Skipping job due to unrecognized job type: {job.get('EmploymentType')}")
-            return None
+        location = Location(
+            city=job.get("City"), state=job.get("State"), country=Country.US_CANADA
+        )
+        job_type = ZipRecruiterScraper.get_job_type_enum(
+            job.get("EmploymentType", "").replace("-", "").lower()
+        )

        formatted_salary = job.get("FormattedSalaryShort", "")
        salary_parts = formatted_salary.split(" ")

        min_salary_str = salary_parts[0][1:].replace(",", "")
-        if '.' in min_salary_str:
+        if "." in min_salary_str:
            min_amount = int(float(min_salary_str) * 1000)
        else:
            min_amount = int(min_salary_str.replace("K", "000"))

        if len(salary_parts) >= 3 and salary_parts[2].startswith("$"):
            max_salary_str = salary_parts[2][1:].replace(",", "")
-            if '.' in max_salary_str:
+            if "." in max_salary_str:
                max_amount = int(float(max_salary_str) * 1000)
            else:
                max_amount = int(max_salary_str.replace("K", "000"))
@@ -232,17 +259,19 @@ class ZipRecruiterScraper(Scraper):
        compensation = Compensation(
            interval=CompensationInterval.YEARLY,
            min_amount=min_amount,
-            max_amount=max_amount
+            max_amount=max_amount,
+            currency="USD/CAD",
        )
        save_job_url = job.get("SaveJobURL", "")
-        posted_time_match = re.search(r"posted_time=(\d{4}-\d{2}-\d{2}T\d{2}:\d{2}:\d{2}Z)", save_job_url)
+        posted_time_match = re.search(
+            r"posted_time=(\d{4}-\d{2}-\d{2}T\d{2}:\d{2}:\d{2}Z)", save_job_url
+        )
        if posted_time_match:
            date_time_str = posted_time_match.group(1)
            date_posted_obj = datetime.strptime(date_time_str, "%Y-%m-%dT%H:%M:%SZ")
            date_posted = date_posted_obj.date()
        else:
            date_posted = date.today()
-        job_url = job.get("JobURL")

        return JobPost(
            title=title,
@@ -257,32 +286,31 @@ class ZipRecruiterScraper(Scraper):
        return job_post

    @staticmethod
-    def job_type_from_string(value: str) -> Optional[JobType]:
-        if not value:
-            return None
+    def get_job_type_enum(job_type_str: str) -> Optional[JobType]:
+        for job_type in JobType:
+            if job_type_str in job_type.value:
+                a = True
+                return job_type
+        return None

-        if value.lower() == "contractor":
-            value = "contract"
-        normalized_value = value.replace("_", "")
-        for item in JobType:
-            if item.value == normalized_value:
-                return item
-        raise ValueError(f"Invalid value for JobType: {value}")
-
-    def get_description(
-            self,
-        job_page_url: str
-    ) -> Tuple[Optional[str], Optional[str]]:
+    def get_description(self, job_page_url: str) -> Tuple[Optional[str], Optional[str]]:
        """
        Retrieves job description by going to the job page url
        :param job_page_url:
        :param session:
        :return: description or None, response url
        """
-        response = self.session.get(
-            job_page_url, headers=ZipRecruiterScraper.headers(), allow_redirects=True
-        )
-        if response.status_code not in range(200, 400):
+        try:
+            response = requests.get(
+                job_page_url,
+                headers=ZipRecruiterScraper.headers(),
+                allow_redirects=True,
+                timeout=5,
+                proxies=self.proxy,
+            )
+            if response.status_code not in range(200, 400):
+                return None, None
+        except Exception as e:
            return None, None

        html_string = response.content
@@ -293,6 +321,36 @@ class ZipRecruiterScraper(Scraper):
            return job_description_div.text.strip(), response.url
        return None, response.url

+    @staticmethod
+    def add_params(scraper_input, page) -> dict[str, str | Any]:
+        params = {
+            "search": scraper_input.search_term,
+            "location": scraper_input.location,
+            "page": page,
+            "form": "jobs-landing",
+        }
+        job_type_value = None
+        if scraper_input.job_type:
+            if scraper_input.job_type.value == "fulltime":
+                job_type_value = "full_time"
+            elif scraper_input.job_type.value == "parttime":
+                job_type_value = "part_time"
+            else:
+                job_type_value = scraper_input.job_type.value
+
+        if job_type_value:
+            params[
+                "refine_by_employment"
+            ] = f"employment_type:employment_type:{job_type_value}"
+
+        if scraper_input.is_remote:
+            params["refine_by_location_type"] = "only_remote"
+
+        if scraper_input.distance:
+            params["radius"] = scraper_input.distance
+
+        return params
+
    @staticmethod
    def get_interval(interval_str: str):
        """
@@ -309,7 +367,7 @@ class ZipRecruiterScraper(Scraper):
        return CompensationInterval(interval_str)

    @staticmethod
-    def get_date_posted(job: BeautifulSoup) -> Optional[datetime.date]:
+    def get_date_posted(job: Tag) -> Optional[datetime.date]:
        """
        Extracts the date a job was posted
        :param job
@@ -335,7 +393,7 @@ class ZipRecruiterScraper(Scraper):
        return None

    @staticmethod
-    def get_compensation(job: BeautifulSoup) -> Optional[Compensation]:
+    def get_compensation(job: Tag) -> Optional[Compensation]:
        """
        Parses the compensation tag from the job BeautifulSoup object
        :param job
@@ -365,7 +423,10 @@ class ZipRecruiterScraper(Scraper):
                amounts.append(amount)

            compensation = Compensation(
-                interval=interval, min_amount=min(amounts), max_amount=max(amounts)
+                interval=interval,
+                min_amount=min(amounts),
+                max_amount=max(amounts),
+                currency="USD/CAD",
            )

            return compensation
@@ -373,7 +434,7 @@ class ZipRecruiterScraper(Scraper):
        return create_compensation_object(pay)

    @staticmethod
-    def get_location(job: BeautifulSoup) -> Location:
+    def get_location(job: Tag) -> Location:
        """
        Extracts the job location from BeatifulSoup object
        :param job:
@@ -389,10 +450,7 @@ class ZipRecruiterScraper(Scraper):
                city, state = None, None
        else:
            city, state = None, None
-        return Location(
-            city=city,
-            state=state,
-        )
+        return Location(city=city, state=state, country=Country.US_CANADA)

    @staticmethod
    def headers() -> dict:
@@ -403,3 +461,9 @@ class ZipRecruiterScraper(Scraper):
        return {
            "User-Agent": "Mozilla/5.0 (Macintosh; Intel Mac OS X 10_14_6) AppleWebKit/537.36 (KHTML, like Gecko) Chrome/78.0.3904.97 Safari/537.36"
        }
+
+    @staticmethod
+    def cleanurl(url):
+        parsed_url = urlparse(url)
+
+        return urlunparse((parsed_url.scheme, parsed_url.netloc, parsed_url.path, parsed_url.params, '', ''))
--- a/src/jobspy/tests/test_indeed.py
+++ b/src/jobspy/tests/test_indeed.py
@@ -1,9 +0,0 @@
-from jobspy import scrape_jobs
-
-
-def test_indeed():
-    result = scrape_jobs(
-        site_name="indeed",
-        search_term="software engineer",
-    )
-    assert result is not None
--- a/src/jobspy/tests/test_linkedin.py
+++ b/src/jobspy/tests/test_linkedin.py
@@ -1,9 +0,0 @@
-from jobspy import scrape_jobs
-
-
-def test_linkedin():
-    result = scrape_jobs(
-        site_name="linkedin",
-        search_term="software engineer",
-    )
-    assert result is not None
--- a/src/jobspy/tests/test_ziprecruiter.py
+++ b/src/jobspy/tests/test_ziprecruiter.py
@@ -1,10 +0,0 @@
-from jobspy import scrape_jobs
-
-
-def test_ziprecruiter():
-    result = scrape_jobs(
-        site_name="zip_recruiter",
-        search_term="software engineer",
-    )
-
-    assert result is not None
--- a/src/jobspy/tests/init.py
+++ b/src/jobspy/tests/init.py
--- a/src/tests/test_all.py
+++ b/src/tests/test_all.py
@@ -0,0 +1,12 @@
+from ..jobspy import scrape_jobs
+import pandas as pd
+
+
+def test_all():
+    result = scrape_jobs(
+        site_name=["linkedin", "indeed", "zip_recruiter"],
+        search_term="software engineer",
+        results_wanted=5,
+    )
+
+    assert isinstance(result, pd.DataFrame) and not result.empty, "Result should be a non-empty DataFrame"
--- a/src/tests/test_indeed.py
+++ b/src/tests/test_indeed.py
@@ -0,0 +1,10 @@
+from ..jobspy import scrape_jobs
+import pandas as pd
+
+
+def test_indeed():
+    result = scrape_jobs(
+        site_name="indeed",
+        search_term="software engineer",
+    )
+    assert isinstance(result, pd.DataFrame) and not result.empty, "Result should be a non-empty DataFrame"
--- a/src/tests/test_linkedin.py
+++ b/src/tests/test_linkedin.py
@@ -0,0 +1,10 @@
+from ..jobspy import scrape_jobs
+import pandas as pd
+
+
+def test_linkedin():
+    result = scrape_jobs(
+        site_name="linkedin",
+        search_term="software engineer",
+    )
+    assert isinstance(result, pd.DataFrame) and not result.empty, "Result should be a non-empty DataFrame"
--- a/src/tests/test_ziprecruiter.py
+++ b/src/tests/test_ziprecruiter.py
@@ -0,0 +1,11 @@
+from ..jobspy import scrape_jobs
+import pandas as pd
+
+
+def test_ziprecruiter():
+    result = scrape_jobs(
+        site_name="zip_recruiter",
+        search_term="software engineer",
+    )
+
+    assert isinstance(result, pd.DataFrame) and not result.empty, "Result should be a non-empty DataFrame"
Author	SHA1	Message	Date
Cullen Watson	af07c1ecbd	add offset param & email extraction (#51 ) * add offset param * [enh]: extract emails	2023-09-28 18:11:28 -05:00
Cullen Watson	286b9e1256	chore: version number	2023-09-21 20:28:57 -05:00
Cullen Watson	162dd40b0f	docs: add usejobspy.com	2023-09-21 20:27:04 -05:00
Cullen Watson	558e352939	fix: job type param bug	2023-09-21 17:42:24 -05:00
Zachary Hampton	efad1a1b7d	Update README.md	2023-09-21 09:52:18 -07:00
Cullen Watson	eaa481c2f4	docs: add macos catalina to faq	2023-09-19 12:50:14 -05:00
Zachary Hampton	b914aa6449	Update README.md	2023-09-16 13:52:30 -07:00
Zachary Hampton	6adbfb8b29	Update README.md	2023-09-16 13:51:45 -07:00
Zachary Hampton	a3b9dd50ff	(docs) homepage	2023-09-15 16:14:26 -07:00
Zachary Hampton	d3ba3a4878	docs: sales call	2023-09-15 11:51:22 -07:00
Cullen Watson	f524789d74	docs: grammar readme	2023-09-15 10:18:24 -05:00
Cullen Watson	f3890d4830	docs: update	2023-09-09 10:55:33 -05:00
Cullen Watson	60c9728691	docs: typo	2023-09-08 12:27:49 -05:00
Cullen Watson	f79d975e5f	docs: clarify - README.md	2023-09-07 13:46:14 -05:00
Cullen Watson	d6368f909b	docs: typo	2023-09-07 13:39:56 -05:00
Cullen Watson	6fcf7f666e	docs: update typo in example	2023-09-07 13:37:53 -05:00
Cullen Watson	4406f9350f	docs: update vid	2023-09-07 13:35:10 -05:00
Cullen Watson	ca5155f234	docs: add feature	2023-09-07 11:36:16 -05:00
Cullen Watson	822a55783e	docs: temp update	2023-09-07 11:35:14 -05:00
Cullen Watson	59f739018a	Proxy support (#44 ) * add proxy support * return as data frame	2023-09-07 11:28:17 -05:00
Zachary Hampton	a37e7f235e	Merge pull request #42 from cullenwatson/fix/class-type-error - refactor & #41 bug fix	2023-09-06 16:33:59 -07:00
Zachary Hampton	690739e858	- refactor & #41 bug fix	2023-09-06 16:32:51 -07:00
Cullen Watson	43eb2fe0e8	remove gitattr	2023-09-06 11:34:51 -05:00
Cullen Watson	e50227bba6	clear output jupyter	2023-09-06 11:32:32 -05:00
Cullen Watson	45c2d76e15	add yt guide	2023-09-06 11:26:55 -05:00
Cullen Watson	fd883178be	Thread sites (#40 )	2023-09-06 09:47:11 -05:00
Cullen Watson	70e2218c67	reduce size of jupyter notebook	2023-09-05 13:09:18 -05:00
Cullen Watson	d6947ecdd7	Update README.md	2023-09-05 13:03:32 -05:00
Cullen Watson	5191658562	Update README.md	2023-09-05 12:27:00 -05:00
Cullen Watson	1c264b8c58	Indeed country support (#38 )	2023-09-05 12:17:22 -05:00
Cullen Watson	1598d4ff63	update README.md	2023-09-04 22:58:46 -05:00
Cullen Watson	bf2460684b	update README.md	2023-09-04 22:52:21 -05:00
Cullen Watson	f5b1e95e64	version number	2023-09-03 20:05:54 -05:00
Cullen Watson	7ae7ecdee8	Validation error (#35 )	2023-09-03 20:05:31 -05:00
Cullen Watson	69b47a2053	docs: clean README.md	2023-09-03 18:13:24 -05:00
Cullen Watson	e7e4acd69c	Update README.md	2023-09-03 18:11:46 -05:00
Cullen Watson	94efc3a099	clean README.md	2023-09-03 18:11:18 -05:00
Cullen Watson	17586dcc28	Update README.md	2023-09-03 18:02:43 -05:00
Cullen Watson	7cc8f4864c	move /tests to /src	2023-09-03 15:40:44 -05:00